@misc{indiciae0c13a64f4f54, title = {StoryTeller: Improving Long Video Description through Global Audio-Visual Character Identification}, author = {Yichen He and Yuan Lin and Jianchao Wu and Hanchong Zhang and Yuchen Zhang and Ruicheng Le}, year = {2025}, url = {https://arxiv.org/abs/2411.07076}, note = {Source identifier: 2411.07076} }