@misc{indiciae2065e78d8996, title = {Cross-modal supervised learning for better acoustic representations}, author = {Shaoyong Jia and Xin Shu and Yang Yang and Dawei Liang and Qiyue Liu and Junhui Liu}, year = {2020}, url = {https://arxiv.org/abs/1911.07917}, note = {Source identifier: 1911.07917} }