@misc{indiciaeebf6b6cb9a37, title = {Lightweight and High-Fidelity End-to-End Text-to-Speech with Multi-Band Generation and Inverse Short-Time Fourier Transform}, author = {Masaya Kawamura and Yuma Shirahata and Ryuichi Yamamoto and Kentaro Tachibana}, year = {2023}, url = {https://arxiv.org/abs/2210.15975}, note = {Source identifier: 2210.15975} }