@misc{indiciae1c10c29977fe, title = {SpeakingFaces: A Large-Scale Multimodal Dataset of Voice Commands with Visual and Thermal Video Streams}, author = {Madina Abdrakhmanova and Askat Kuzdeuov and Sheikh Jarju and Yerbolat Khassanov and Michael Lewis and Huseyin Atakan Varol}, year = {2021}, url = {https://arxiv.org/abs/2012.02961}, note = {Source identifier: 2012.02961} }