@misc{indiciae4c2191bcb1a3, title = {Vision-Language Models for Autonomous Driving: CLIP-Based Dynamic Scene Understanding}, author = {Mohammed Elhenawy and Huthaifa I. Ashqar and Andry Rakotonirainy and Taqwa I. Alhadidi and Ahmed Jaber and Mohammad Abu Tami}, year = {2025}, url = {https://arxiv.org/abs/2501.05566}, note = {Source identifier: 2501.05566} }