@misc{indiciae8af9f852f55b, title = {DocSLM: A Small Vision-Language Model for Long Multimodal Document Understanding}, author = {Tanveer Hannan and Dimitrios Mallios and Parth Pathak and Faegheh Sardari and Thomas Seidl and Gedas Bertasius and Mohsen Fayyaz and Sunando Sengupta}, year = {2025}, url = {https://arxiv.org/abs/2511.11313}, note = {Source identifier: 2511.11313} }