@misc{indiciae26113473489a, title = {Seeing voices and hearing voices: learning discriminative embeddings using cross-modal self-supervision}, author = {Soo-Whan Chung and Hong Goo Kang and Joon Son Chung}, year = {2020}, doi = {10.21437/interspeech.2020-1113}, url = {https://arxiv.org/abs/2004.14326}, note = {Source identifier: 2004.14326} }