@misc{indiciae516bcee41d16, title = {AudioToken: Adaptation of Text-Conditioned Diffusion Models for Audio-to-Image Generation}, author = {Guy Yariv and Itai Gat and Lior Wolf and Yossi Adi and Idan Schwartz}, year = {2023}, url = {https://arxiv.org/abs/2305.13050}, note = {Source identifier: 2305.13050} }