@misc{indiciae3454e1adeab3, title = {Audio-visual fine-tuning of audio-only ASR models}, author = {Avner May and Dmitriy Serdyuk and Ankit Parag Shah and Otavio Braga and Olivier Siohan}, year = {2023}, url = {https://arxiv.org/abs/2312.09369}, note = {Source identifier: 2312.09369} }