@misc{indiciaefd9f1c0b7833, title = {Single-stage TTS with Masked Audio Token Modeling and Semantic Knowledge Distillation}, author = {Gerard I. Gállego and Roy Fejgin and Chunghsin Yeh and Xiaoyu Liu and Gautam Bhattacharya}, year = {2024}, url = {https://arxiv.org/abs/2409.11003}, note = {Source identifier: 2409.11003} }