@misc{indiciae8757f7096d45, title = {SurgLLM: A Versatile Large Multimodal Model with Spatial Focus and Temporal Awareness for Surgical Video Understanding}, author = {Zhen Chen and Xingjian Luo and Kun Yuan and Jinlin Wu and Danny T. M. Chan and Nassir Navab and Hongbin Liu and Zhen Lei and Jiebo Luo}, year = {2025}, url = {https://arxiv.org/abs/2509.00357}, note = {Source identifier: 2509.00357} }