@misc{indiciaeb2b3ab01ad0c, title = {Interpreting Video Representations with Spatio-Temporal Sparse Autoencoders}, author = {Atahan Dokme and Sriram Vishwanath}, year = {2026}, doi = {10.1145/3767308.3836082}, url = {https://arxiv.org/abs/2604.03919}, note = {Source identifier: 2604.03919} }