@misc{indiciaed335eb4835cf, title = {Efficient Vocabulary-Free Fine-Grained Visual Recognition in the Age of Multimodal LLMs}, author = {Hari Chandana Kuchibhotla and Sai Srinivas Kancheti and Abbavaram Gowtham Reddy and Vineeth N Balasubramanian}, year = {2025}, url = {https://arxiv.org/abs/2505.01064}, note = {Source identifier: 2505.01064} }