@misc{indiciaecd31edf08c9f, title = {LUT-LLM: Efficient Large Language Model Inference with Memory-based Computations on FPGAs}, author = {Zifan He and Shengyu Ye and Rui Ma and Yang Wang and Jason Cong}, year = {2026}, url = {https://arxiv.org/abs/2511.06174}, note = {Source identifier: 2511.06174} }