@misc{indiciaee0573ee9a931, title = {DOA-Aware Audio-Visual Self-Supervised Learning for Sound Event Localization and Detection}, author = {Yoto Fujita and Yoshiaki Bando and Keisuke Imoto and Masaki Onishi and Kazuyoshi Yoshii}, year = {2024}, url = {https://arxiv.org/abs/2410.22803}, note = {Source identifier: 2410.22803} }