@misc{indiciaea3ec1aefd408, title = {PaCa-ViT: Learning Patch-to-Cluster Attention in Vision Transformers}, author = {Ryan Grainger and Thomas Paniagua and Xi Song and Naresh Cuntoor and Mun Wai Lee and Tianfu Wu}, year = {2023}, url = {https://arxiv.org/abs/2203.11987}, note = {Source identifier: 2203.11987} }