@misc{indiciaeb3e568794b42, title = {ZMM-TTS: Zero-shot Multilingual and Multispeaker Speech Synthesis Conditioned on Self-supervised Discrete Speech Representations}, author = {Cheng Gong and Xin Wang and Erica Cooper and Dan Wells and Longbiao Wang and Jianwu Dang and Korin Richmond and Junichi Yamagishi}, year = {2024}, url = {https://arxiv.org/abs/2312.14398}, note = {Source identifier: 2312.14398} }