@misc{indiciae323bdc4c2a2d, title = {ETTA: Elucidating the Design Space of Text-to-Audio Models}, author = {Sang-gil Lee and Zhifeng Kong and Arushi Goel and Sungwon Kim and Rafael Valle and Bryan Catanzaro}, year = {2025}, url = {https://arxiv.org/abs/2412.19351}, note = {Source identifier: 2412.19351} }