@misc{indiciaeaaef5d3787a8, title = {An Overview of Deep-Learning-Based Audio-Visual Speech Enhancement and Separation}, author = {Daniel Michelsanti and Zheng-Hua Tan and Shi-Xiong Zhang and Yong Xu and Meng Yu and Dong Yu and Jesper Jensen}, year = {2021}, url = {https://arxiv.org/abs/2008.09586}, note = {Source identifier: 2008.09586} }