@misc{indiciae2e8c910d01d9, title = {Audio-Visual Speaker Diarization Based on Spatiotemporal Bayesian Fusion}, author = {Israel D. Gebru and Silèye Ba and Xiaofei Li and Radu Horaud}, year = {2016}, doi = {10.1109/tpami.2017.2648793}, url = {https://arxiv.org/abs/1603.09725}, note = {Source identifier: 1603.09725} }