@misc{indiciaebfaeb31b293d, title = {Auffusion: Leveraging the Power of Diffusion and Large Language Models for Text-to-Audio Generation}, author = {Jinlong Xue and Yayue Deng and Yingming Gao and Ya Li}, year = {2024}, url = {https://arxiv.org/abs/2401.01044}, note = {Source identifier: 2401.01044} }