@misc{indiciae7359b73a3f9a, title = {Total-Duration-Aware Duration Modeling for Text-to-Speech Systems}, author = {Sefik Emre Eskimez and Xiaofei Wang and Manthan Thakker and Chung-Hsien Tsai and Canrun Li and Zhen Xiao and Hemin Yang and Zirun Zhu and Min Tang and Jinyu Li and Sheng Zhao and Naoyuki Kanda}, year = {2024}, url = {https://arxiv.org/abs/2406.04281}, note = {Source identifier: 2406.04281} }