@misc{indiciaee35186330b1e, title = {Depth-Wise Representation Development Under Blockwise Self-Supervised Learning for Video Vision Transformers}, author = {Jonas Römer and Timo Dickscheid}, year = {2026}, url = {https://arxiv.org/abs/2601.09040}, note = {Source identifier: 2601.09040} }