@misc{indiciaeaa0c893c3284, title = {Deep Triplet Neural Networks with Cluster-CCA for Audio-Visual Cross-modal Retrieval}, author = {Donghuo Zeng and Yi Yu and Keizo Oyama}, year = {2021}, url = {https://arxiv.org/abs/1908.03737}, note = {Source identifier: 1908.03737} }