@misc{indiciae890fd1cc313c, title = {SoundNet: Learning Sound Representations from Unlabeled Video}, author = {Yusuf Aytar and Carl Vondrick and Antonio Torralba}, year = {2016}, url = {https://arxiv.org/abs/1610.09001}, note = {Source identifier: 1610.09001} }