@misc{indiciaeab7a9df5bb73, title = {JDI-T: Jointly trained Duration Informed Transformer for Text-To-Speech without Explicit Alignment}, author = {Dan Lim and Won Jang and Gyeonghwan O and Heayoung Park and Bongwan Kim and Jaesam Yoon}, year = {2020}, url = {https://arxiv.org/abs/2005.07799}, note = {Source identifier: 2005.07799} }