@misc{indiciaed601f3c930e9, title = {Visual-Word Tokenizer: Beyond Fixed Sets of Tokens in Vision Transformers}, author = {Leonidas Gee and Wing Yan Li and Viktoriia Sharmanska and Novi Quadrianto}, year = {2025}, url = {https://arxiv.org/abs/2411.15397}, note = {Source identifier: 2411.15397} }