@misc{indiciae5c27f83aa583, title = {Attentive Fusion Enhanced Audio-Visual Encoding for Transformer Based Robust Speech Recognition}, author = {Liangfa Wei and Jie Zhang and Junfeng Hou and Lirong Dai}, year = {2020}, url = {https://arxiv.org/abs/2008.02686}, note = {Source identifier: 2008.02686} }