@misc{indiciaee213e4f010e8, title = {DMP-TTS: Disentangled multi-modal Prompting for Controllable Text-to-Speech with Chained Guidance}, author = {Kang Yin and Chunyu Qiang and Sirui Zhao and Xiaopeng Wang and Yuzhe Liang and Pengfei Cai and Tong Xu and Chen Zhang and Enhong Chen}, year = {2025}, url = {https://arxiv.org/abs/2512.09504}, note = {Source identifier: 2512.09504} }