@misc{indiciae27031ed0b686, title = {Fine-Tuning Text-to-Speech Diffusion Models Using Reinforcement Learning with Human Feedback}, author = {Jingyi Chen and Ju Seung Byun and Micha Elsner and Pichao Wang and Andrew Perrault}, year = {2025}, url = {https://arxiv.org/abs/2508.03123}, note = {Source identifier: 2508.03123} }