@misc{indiciae85ea77eec550, title = {Personalized Video Summarization by Multimodal Video Understanding}, author = {Brian Chen and Xiangyuan Zhao and Yingnan Zhu}, year = {2024}, doi = {10.1145/3627673.3680011}, url = {https://arxiv.org/abs/2411.03531}, note = {Source identifier: 2411.03531} }