@misc{indiciaed73dcbcb9445, title = {Bridging the Visual Gap: Fine-Tuning Multimodal Models with Knowledge-Adapted Captions}, author = {Moran Yanuka and Assaf Ben Kish and Yonatan Bitton and Idan Szpektor and Raja Giryes}, year = {2025}, doi = {10.18653/v1/2025.naacl-long.527}, url = {https://arxiv.org/abs/2411.09018}, note = {Source identifier: 2411.09018} }