@misc{indiciae027ed46268e8, title = {MobileCLIP: Fast Image-Text Models through Multi-Modal Reinforced Training}, author = {Pavan Kumar Anasosalu Vasu and Hadi Pouransari and Fartash Faghri and Raviteja Vemulapalli and Oncel Tuzel}, year = {2024}, url = {https://arxiv.org/abs/2311.17049}, note = {Source identifier: 2311.17049} }