@misc{indiciaeabe0e462b334, title = {AV2AV: Direct Audio-Visual Speech to Audio-Visual Speech Translation with Unified Audio-Visual Speech Representation}, author = {Jeongsoo Choi and Se Jin Park and Minsu Kim and Yong Man Ro}, year = {2024}, url = {https://arxiv.org/abs/2312.02512}, note = {Source identifier: 2312.02512} }