@misc{indiciaedbdba6d07b3b, title = {Fast Text-to-Audio Generation with One-Step Sampling via Energy-Scoring and Auxiliary Contextual Representation Distillation}, author = {Kuan-Po Huang and Bo-Ru Lu and Byeonggeun Kim and Mihee Lee and Zalan Fabian and Renard Korzeniowski and Qingming Tang and Greg Ver Steeg and Hung-yi Lee and Chieh-Chi Kao and Chao Wang}, year = {2026}, url = {https://arxiv.org/abs/2605.00329}, note = {Source identifier: 2605.00329} }