@misc{indiciaed1a76a973c5b, title = {Self-Supervised Learning of Audio-Visual Objects from Video}, author = {Triantafyllos Afouras and Andrew Owens and Joon Son Chung and Andrew Zisserman}, year = {2020}, url = {https://arxiv.org/abs/2008.04237}, note = {Source identifier: 2008.04237} }