@misc{indiciae6a294d80c203, title = {VITS2: Improving Quality and Efficiency of Single-Stage Text-to-Speech with Adversarial Learning and Architecture Design}, author = {Jungil Kong and Jihoon Park and Beomjeong Kim and Jeongmin Kim and Dohee Kong and Sangjin Kim}, year = {2023}, url = {https://arxiv.org/abs/2307.16430}, note = {Source identifier: 2307.16430} }