@misc{indiciae1b4e99ff1400, title = {VideoSEMA: a scalable and efficient Mamba-like attention for video understanding}, author = {Nhat Thanh Tran and Fanghui Xue and Shuai Zhang and Jiancheng Lyu and Yunling Zheng and Yingyong Qi and Jack Xin}, year = {2026}, url = {https://arxiv.org/abs/2607.14711}, note = {Source identifier: 2607.14711} }