@article{violetavisionlanguagemodelforarabic, title = {Violet: A Vision-Language Model for Arabic Image Captioning with Gemini Decoder}, author = {Abdelrahman Mohamed and Fakhraddin Alwajih and El Moatez Billah Nagoudi and Alcides Alcoba Inciarte and Muhammad Abdul-Mageed}, year = {2023}, eprint = {2311.08844}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2311.08844v1}, }