@misc{indiciaefef25f68499f, title = {Improving Text-To-Audio Models with Synthetic Captions}, author = {Zhifeng Kong and Sang-gil Lee and Deepanway Ghosal and Navonil Majumder and Ambuj Mehrish and Rafael Valle and Soujanya Poria and Bryan Catanzaro}, year = {2024}, url = {https://arxiv.org/abs/2406.15487}, note = {Source identifier: 2406.15487} }