@misc{indiciaedcb1dfc09a51, title = {MINI-LLM: Memory-Efficient Structured Pruning for Large Language Models}, author = {Hongrong Cheng and Miao Zhang and Javen Qinfeng Shi}, year = {2024}, url = {https://arxiv.org/abs/2407.11681}, note = {Source identifier: 2407.11681} }