@misc{indiciae06d7ea38b5d1, title = {FastVLM: Efficient Vision Encoding for Vision Language Models}, author = {Pavan Kumar Anasosalu Vasu and Fartash Faghri and Chun-Liang Li and Cem Koc and Nate True and Albert Antony and Gokul Santhanam and James Gabriel and Peter Grasch and Oncel Tuzel and Hadi Pouransari}, year = {2025}, url = {https://arxiv.org/abs/2412.13303}, note = {Source identifier: 2412.13303} }