@misc{indiciae4d263483d564, title = {When Vision Models Meet Parameter Efficient Look-Aside Adapters Without Large-Scale Audio Pretraining}, author = {Juan Yeo and Jinkwan Jang and Kyubyung Chae and Seongkyu Mun and Taesup Kim}, year = {2024}, url = {https://arxiv.org/abs/2412.05951}, note = {Source identifier: 2412.05951} }