@misc{indiciae3dec1d60ea82, title = {Mellotron: Multispeaker expressive voice synthesis by conditioning on rhythm, pitch and global style tokens}, author = {Rafael Valle and Jason Li and Ryan Prenger and Bryan Catanzaro}, year = {2019}, url = {https://arxiv.org/abs/1910.11997}, note = {Source identifier: 1910.11997} }