@misc{indiciaebd2084a55e93, title = {Fusing information streams in end-to-end audio-visual speech recognition}, author = {Wentao Yu and Steffen Zeiler and Dorothea Kolossa}, year = {2021}, url = {https://arxiv.org/abs/2104.09482}, note = {Source identifier: 2104.09482} }