@misc{indiciaea7df8c36c636, title = {End-To-End Audiovisual Feature Fusion for Active Speaker Detection}, author = {Fiseha B. Tesema and Zheyuan Lin and Shiqiang Zhu and Wei Song and Jason Gu and Hong Wu}, year = {2022}, doi = {10.1117/12.2643881}, url = {https://arxiv.org/abs/2207.13434}, note = {Source identifier: 2207.13434} }