@misc{indiciae0bfa685f8fe1, title = {Learning Speech Representations from Raw Audio by Joint Audiovisual Self-Supervision}, author = {Abhinav Shukla and Stavros Petridis and Maja Pantic}, year = {2020}, url = {https://arxiv.org/abs/2007.04134}, note = {Source identifier: 2007.04134} }