@misc{indiciae57d7bd306b4c, title = {Putting a Face to the Voice: Fusing Audio and Visual Signals Across a Video to Determine Speakers}, author = {Ken Hoover and Sourish Chaudhuri and Caroline Pantofaru and Malcolm Slaney and Ian Sturdy}, year = {2017}, url = {https://arxiv.org/abs/1706.00079}, note = {Source identifier: 1706.00079} }