@misc{indiciae8a9fa9b86f77, title = {SimpleSpeech 2: Towards Simple and Efficient Text-to-Speech with Flow-based Scalar Latent Transformer Diffusion Models}, author = {Dongchao Yang and Rongjie Huang and Yuanyuan Wang and Haohan Guo and Dading Chong and Songxiang Liu and Xixin Wu and Helen Meng}, year = {2024}, url = {https://arxiv.org/abs/2408.13893}, note = {Source identifier: 2408.13893} }