@misc{indiciaec52014b7a929, title = {Transduce and Speak: Neural Transducer for Text-to-Speech with Semantic Token Prediction}, author = {Minchan Kim and Myeonghun Jeong and Byoung Jin Choi and Dongjune Lee and Nam Soo Kim}, year = {2023}, url = {https://arxiv.org/abs/2311.02898}, note = {Source identifier: 2311.02898} }