@article{visionlanguagemodelsforautonomousdriving, title = {Vision-Language Models for Autonomous Driving: CLIP-Based Dynamic Scene Understanding}, author = {Mohammed Elhenawy and Huthaifa I. Ashqar and Andry Rakotonirainy and Taqwa I. Alhadidi and Ahmed Jaber and Mohammad Abu Tami}, year = {2025}, eprint = {2501.05566}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2501.05566v1}, }