@misc{indiciae40cf89a5243d, title = {Learning Lip-Based Audio-Visual Speaker Embeddings with AV-HuBERT}, author = {Bowen Shi and Abdelrahman Mohamed and Wei-Ning Hsu}, year = {2022}, url = {https://arxiv.org/abs/2205.07180}, note = {Source identifier: 2205.07180} }