@misc{indiciaeb12265aff3f6, title = {Distilling Vision-Language Pretraining for Efficient Cross-Modal Retrieval}, author = {Young Kyun Jang and Donghyun Kim and Ser-nam Lim}, year = {2024}, url = {https://arxiv.org/abs/2405.14726}, note = {Source identifier: 2405.14726} }