@misc{indiciaeda77dde1fc96, title = {How Vision-Language Tasks Benefit from Large Pre-trained Models: A Survey}, author = {Yayun Qi and Hongxi Li and Yiqi Song and Xinxiao Wu and Jiebo Luo}, year = {2024}, url = {https://arxiv.org/abs/2412.08158}, note = {Source identifier: 2412.08158} }