@misc{indiciae9198bcb6dea6, title = {End-to-End Multi-Person Audio/Visual Automatic Speech Recognition}, author = {Otavio Braga and Takaki Makino and Olivier Siohan and Hank Liao}, year = {2022}, url = {https://arxiv.org/abs/2205.05586}, note = {Source identifier: 2205.05586} }