@misc{indiciae8070fc3d0b46, title = {Summarization of Multimodal Presentations with Vision-Language Models: Study of the Effect of Modalities and Structure}, author = {Théo Gigant and Camille Guinaudeau and Frédéric Dufaux}, year = {2025}, url = {https://arxiv.org/abs/2504.10049}, note = {Source identifier: 2504.10049} }