@misc{indiciaef888406f2bbe, title = {Learning to Mask and Permute Visual Tokens for Vision Transformer Pre-Training}, author = {Lorenzo Baraldi and Roberto Amoroso and Marcella Cornia and Andrea Pilzer and Rita Cucchiara}, year = {2025}, url = {https://arxiv.org/abs/2306.07346}, note = {Source identifier: 2306.07346} }