@misc{indiciaefceef0fb15bc, title = {FAM-HRI: Foundation-Model Assisted Multi-Modal Human-Robot Interaction Combining Gaze and Speech}, author = {Yuzhi Lai and Shenghai Yuan and Peizheng Li and Boya Zhang and Benjamin Kiefer and Tianchen Deng and Andreas Zell}, year = {2026}, url = {https://arxiv.org/abs/2503.16492}, note = {Source identifier: 2503.16492} }