@misc{indiciae2269b4d0b94b, title = {Zero-shot text-to-speech synthesis conditioned using self-supervised speech representation model}, author = {Kenichi Fujita and Takanori Ashihara and Hiroki Kanagawa and Takafumi Moriya and Yusuke Ijima}, year = {2023}, doi = {10.1109/icasspw59220.2023.10193459}, url = {https://arxiv.org/abs/2304.11976}, note = {Source identifier: 2304.11976} }