@misc{indiciaec6aeee600ba7, title = {Cross-modal Prompts: Adapting Large Pre-trained Models for Audio-Visual Downstream Tasks}, author = {Haoyi Duan and Yan Xia and Mingze Zhou and Li Tang and Jieming Zhu and Zhou Zhao}, year = {2023}, url = {https://arxiv.org/abs/2311.05152}, note = {Source identifier: 2311.05152} }