@misc{indiciae715c6d3054b1, title = {DiffStyleTTS: Diffusion-based Hierarchical Prosody Modeling for Text-to-Speech with Diverse and Controllable Styles}, author = {Jiaxuan Liu and Zhaoci Liu and Yajun Hu and Yingying Gao and Shilei Zhang and Zhenhua Ling}, year = {2024}, url = {https://arxiv.org/abs/2412.03388}, note = {Source identifier: 2412.03388} }