@misc{indiciae531bcddc0376, title = {Leveraging Vision-Language Foundation Models for Fine-Grained Downstream Tasks}, author = {Denis Coquenet and Clément Rambour and Emanuele Dalsasso and Nicolas Thome}, year = {2023}, url = {https://arxiv.org/abs/2307.06795}, note = {Source identifier: 2307.06795} }