@misc{indiciaec2431a8b62c4, title = {SimpleSpeech: Towards Simple and Efficient Text-to-Speech with Scalar Latent Transformer Diffusion Models}, author = {Dongchao Yang and Dingdong Wang and Haohan Guo and Xueyuan Chen and Xixin Wu and Helen Meng}, year = {2024}, url = {https://arxiv.org/abs/2406.02328}, note = {Source identifier: 2406.02328} }