@misc{indiciae6fce8f7d01c7, title = {High Fidelity Text-to-Speech Via Discrete Tokens Using Token Transducer and Group Masked Language Model}, author = {Joun Yeop Lee and Myeonghun Jeong and Minchan Kim and Ji-Hyun Lee and Hoon-Young Cho and Nam Soo Kim}, year = {2024}, url = {https://arxiv.org/abs/2406.17310}, note = {Source identifier: 2406.17310} }