@misc{indiciae90afdf24ca57, title = {Learning Visual-Audio Representations for Voice-Controlled Robots}, author = {Peixin Chang and Shuijing Liu and D. Livingston McPherson and Katherine Driggs-Campbell}, year = {2023}, url = {https://arxiv.org/abs/2109.02823}, note = {Source identifier: 2109.02823} }