<?xml version="1.0" encoding="US-ASCII"?>
<dblp>
<inproceedings key="conf/asplos/HeMGSGLLWM25" mdate="2025-05-09">
<author orcid="0000-0003-3054-0617">Yintao He</author>
<author orcid="0000-0002-7393-4504">Haiyu Mao</author>
<author orcid="0000-0003-0162-4547">Christina Giannoula</author>
<author orcid="0000-0002-4029-0175">Mohammad Sadrosadati</author>
<author orcid="0000-0002-6514-1571">Juan G&#243;mez-Luna</author>
<author orcid="0000-0001-8082-4218">Huawei Li 0001</author>
<author orcid="0000-0002-0874-814X">Xiaowei Li 0001</author>
<author orcid="0000-0001-5172-4736">Ying Wang 0001</author>
<author orcid="0000-0002-0075-2312">Onur Mutlu</author>
<title>PAPI: Exploiting Dynamic Parallelism in Large Language Model Decoding with a Processing-In-Memory-Enabled Computing System.</title>
<pages>766-782</pages>
<year>2025</year>
<booktitle>ASPLOS (2)</booktitle>
<ee>https://doi.org/10.1145/3676641.3716009</ee>
<crossref>conf/asplos/2025-2</crossref>
<url>db/conf/asplos/asplos2025-2.html#HeMGSGLLWM25</url>
</inproceedings>
</dblp>
