@misc{indiciae29a88376dd28, title = {Towards an Automated Multimodal Approach for Video Summarization: Building a Bridge Between Text, Audio and Facial Cue-Based Summarization}, author = {Md Moinul Islam and Sofoklis Kakouros and Janne Heikkilä and Mourad Oussalah}, year = {2025}, url = {https://arxiv.org/abs/2506.23714}, note = {Source identifier: 2506.23714} }