@misc{indiciae1ff0885edee6, title = {Face-StyleSpeech: Enhancing Zero-shot Speech Synthesis from Face Images with Improved Face-to-Speech Mapping}, author = {Minki Kang and Wooseok Han and Eunho Yang}, year = {2024}, url = {https://arxiv.org/abs/2311.05844}, note = {Source identifier: 2311.05844} }