@misc{indiciae45fb4e47c2e3, title = {Robust Wake Word Spotting With Frame-Level Cross-Modal Attention Based Audio-Visual Conformer}, author = {Haoxu Wang and Ming Cheng and Qiang Fu and Ming Li}, year = {2024}, url = {https://arxiv.org/abs/2403.01700}, note = {Source identifier: 2403.01700} }