@misc{indiciaefdc7e98c7f97, title = {V2A-Mapper: A Lightweight Solution for Vision-to-Audio Generation by Connecting Foundation Models}, author = {Heng Wang and Jianbo Ma and Santiago Pascual and Richard Cartwright and Weidong Cai}, year = {2023}, url = {https://arxiv.org/abs/2308.09300}, note = {Source identifier: 2308.09300} }