@misc{indiciae9484969fd2a1, title = {Multimodal Human-Autonomous Agents Interaction Using Pre-Trained Language and Visual Foundation Models}, author = {Linus Nwankwo and Elmar Rueckert}, year = {2024}, url = {https://arxiv.org/abs/2403.12273}, note = {Source identifier: 2403.12273} }