@misc{indiciae3ab2f0aaefaf, title = {Scalable and Explainable Learner-Video Interaction Prediction using Multimodal Large Language Models}, author = {Dominik Glandorf and Fares Fawzi and Tanja Käser}, year = {2026}, url = {https://arxiv.org/abs/2604.04482}, note = {Source identifier: 2604.04482} }