@misc{indiciaeca5e6b239f1b, title = {DoraemonGPT: Toward Understanding Dynamic Scenes with Large Language Models (Exemplified as A Video Agent)}, author = {Zongxin Yang and Guikun Chen and Xiaodi Li and Wenguan Wang and Yi Yang}, year = {2025}, url = {https://arxiv.org/abs/2401.08392}, note = {Source identifier: 2401.08392} }