@misc{indiciae2a79c91e8b7c, title = {NanoVLMs: How small can we go and still make coherent Vision Language Models?}, author = {Mukund Agarwalla and Himanshu Kumar and Raj Dandekar and Rajat Dandekar and Sreedath Panat}, year = {2025}, url = {https://arxiv.org/abs/2502.07838}, note = {Source identifier: 2502.07838} }