<?xml version="1.0"?>
<dblpperson name="Devendar Bureddy" pid="27/10106" n="15">
<person key="homepages/27/10106" mdate="2020-06-16">
<author pid="27/10106">Devendar Bureddy</author>
<url>https://scholar.google.com/citations?user=mWjNqW0AAAAJ</url>
</person>
<r><article key="journals/micro/VenkataPLBALBDS25" mdate="2025-06-11">
<author orcid="0000-0002-5282-1682" pid="55/6076">Manjunath Gorentla Venkata</author>
<author orcid="0009-0009-1551-3048" pid="385/5435">Valentine Petrov</author>
<author orcid="0009-0001-0881-7299" pid="25/9629">Sergey Lebedev</author>
<author orcid="0009-0006-4638-2456" pid="27/10106">Devendar Bureddy</author>
<author orcid="0000-0002-4208-6493" pid="19/7905">Ferrol Aderholdt</author>
<author orcid="0009-0003-7956-1387" pid="70/1942">Joshua Ladd</author>
<author orcid="0009-0004-6224-9802" pid="99/2609">Gil Bloch</author>
<author orcid="0009-0000-4388-7084" pid="129/4246">Mike Dubman</author>
<author orcid="0009-0000-9318-9620" pid="52/6581">Gilad Shainer</author>
<title>Unified Collective Communication: A Unified Library for CPU, GPU, and DPU Collectives.</title>
<pages>26-35</pages>
<year>2025</year>
<month>March - April</month>
<volume>45</volume>
<journal>IEEE Micro</journal>
<number>2</number>
<ee>https://doi.org/10.1109/MM.2025.3534638</ee>
<url>db/journals/micro/micro45.html#VenkataPLBALBDS25</url>
<stream>streams/journals/micro</stream>
</article>
</r>
<r><inproceedings key="conf/hoti/VenkataPLBALBDS24" mdate="2024-09-19">
<author pid="55/6076">Manjunath Gorentla Venkata</author>
<author pid="385/5435">Valentine Petrov</author>
<author pid="25/9629">Sergey Lebedev</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="19/7905">Ferrol Aderholdt</author>
<author pid="70/1942">Joshua Ladd</author>
<author pid="99/2609">Gil Bloch</author>
<author pid="129/4246">Mike Dubman</author>
<author pid="52/6581">Gilad Shainer</author>
<title>Unified Collective Communication (UCC): An Unified Library for CPU, GPU, and DPU Collectives.</title>
<pages>37-46</pages>
<year>2024</year>
<booktitle>HOTI</booktitle>
<ee>https://doi.org/10.1109/HOTI63208.2024.00018</ee>
<crossref>conf/hoti/2024</crossref>
<url>db/conf/hoti/hoti2024.html#VenkataPLBALBDS24</url>
<stream>streams/conf/hoti</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/supercomputer/GrahamLBBSCEKLM20" mdate="2026-06-26">
<author pid="86/5514">Richard L. Graham</author>
<author pid="194/1364">Lion Levi</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="99/2609">Gil Bloch</author>
<author pid="52/6581">Gilad Shainer</author>
<author pid="74/2156">David Cho</author>
<author pid="267/4443">George Elias</author>
<author pid="93/6085">Daniel Klein</author>
<author pid="70/1942">Joshua Ladd</author>
<author pid="267/4666">Ophir Maor</author>
<author pid="267/4264">Ami Marelli</author>
<author pid="76/300">Valentin Petrov</author>
<author pid="267/4272">Evyatar Romlet</author>
<author orcid="0009-0003-1460-3574" pid="20/4298">Yong Qin</author>
<author pid="267/4288">Ido Zemah</author>
<title>Scalable Hierarchical Aggregation and Reduction Protocol (SHARP)<sup>TM</sup> Streaming-Aggregation Hardware Design and Evaluation.</title>
<pages>41-59</pages>
<year>2020</year>
<booktitle>ISC High Performance</booktitle>
<ee type="oa">https://doi.org/10.1007/978-3-030-50743-5_3</ee>
<crossref>conf/supercomputer/2020</crossref>
<url>db/conf/supercomputer/isc2020.html#GrahamLBBSCEKLM20</url>
</inproceedings>
</r>
<r><article key="journals/superfri/GrahamBBSS17" mdate="2020-12-15">
<author pid="86/5514">Richard L. Graham</author>
<author pid="99/2609">Gil Bloch</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="52/6581">Gilad Shainer</author>
<author pid="98/5275">Brian Smith</author>
<title>Towards A Data Centric System Architecture: SHARP.</title>
<pages>4-16</pages>
<year>2017</year>
<volume>4</volume>
<journal>Supercomput. Front. Innov.</journal>
<number>4</number>
<ee type="oa">https://doi.org/10.14529/jsfi170401</ee>
<url>db/journals/superfri/superfri4.html#GrahamBBSS17</url>
</article>
</r>
<r><inproceedings key="conf/sc/GrahamBLRSBGDKK16" mdate="2017-05-23">
<author pid="86/5514">Richard L. Graham</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="47/9777">Pak Lui</author>
<author pid="194/1388">Hal Rosenstock</author>
<author pid="52/6581">Gilad Shainer</author>
<author pid="99/2609">Gil Bloch</author>
<author pid="07/1366">Dror Goldenberg</author>
<author pid="129/4246">Mike Dubman</author>
<author pid="194/1376">Sasha Kotchubievsky</author>
<author pid="194/1355">Vladimir Koushnir</author>
<author pid="194/1364">Lion Levi</author>
<author pid="194/1357">Alex Margolin</author>
<author pid="194/1384">Tamir Ronen</author>
<author pid="09/9457">Alexander Shpiner</author>
<author pid="194/1352">Oded Wertheim</author>
<author pid="72/7892">Eitan Zahavi</author>
<title>Scalable Hierarchical Aggregation Protocol (SHArP): A Hardware Architecture for Efficient Data Reduction.</title>
<pages>1-10</pages>
<year>2016</year>
<booktitle>COMHPC@SC</booktitle>
<ee>https://doi.org/10.1109/COMHPC.2016.006</ee>
<ee>http://dl.acm.org/citation.cfm?id=3018059</ee>
<crossref>conf/sc/2016comhpc</crossref>
<url>db/conf/sc/comhpc2016.html#GrahamBLRSBGDKK16</url>
</inproceedings>
</r>
<r><article key="journals/tpds/WangPBRP14" mdate="2021-10-14">
<author orcid="0000-0003-3557-6301" pid="w/HaoWang-2">Hao Wang 0002</author>
<author pid="69/8165">Sreeram Potluri</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="56/3932">Carlos Rosales</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>GPU-Aware MPI on RDMA-Enabled Clusters: Design, Implementation and Evaluation.</title>
<pages>2595-2605</pages>
<year>2014</year>
<volume>25</volume>
<journal>IEEE Trans. Parallel Distributed Syst.</journal>
<number>10</number>
<ee>https://doi.org/10.1109/TPDS.2013.222</ee>
<ee>http://doi.ieeecomputersociety.org/10.1109/TPDS.2013.222</ee>
<url>db/journals/tpds/tpds25.html#WangPBRP14</url>
</article>
</r>
<r><inproceedings key="conf/ccgrid/PotluriVBKP13" mdate="2023-03-24">
<author pid="69/8165">Sreeram Potluri</author>
<author pid="78/6299">Akshay Venkatesh</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="91/7432">Krishna Chaitanya Kandalla</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>Efficient Intra-node Communication on Intel-MIC Clusters.</title>
<pages>128-135</pages>
<year>2013</year>
<booktitle>CCGRID</booktitle>
<ee>https://doi.org/10.1109/CCGrid.2013.86</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/CCGrid.2013.86</ee>
<crossref>conf/ccgrid/2013</crossref>
<url>db/conf/ccgrid/ccgrid2013.html#PotluriVBKP13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cluster/SubramoniBKSBPAP13" mdate="2023-03-23">
<author pid="65/1958">Hari Subramoni</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="91/7432">Krishna Chaitanya Kandalla</author>
<author pid="45/5668">Karl W. Schulz</author>
<author pid="08/1485">Bill Barth</author>
<author pid="26/2675">Jonathan L. Perkins</author>
<author pid="131/4985">Mark Daniel Arnold</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>Design of network topology aware scheduling services for large InfiniBand clusters.</title>
<pages>1-8</pages>
<year>2013</year>
<booktitle>CLUSTER</booktitle>
<ee>https://doi.org/10.1109/CLUSTER.2013.6702677</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/CLUSTER.2013.6702677</ee>
<crossref>conf/cluster/2013</crossref>
<url>db/conf/cluster/cluster2013.html#SubramoniBKSBPAP13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/hoti/KandallaVHPBP13" mdate="2023-03-24">
<author pid="91/7432">Krishna Chaitanya Kandalla</author>
<author pid="78/6299">Akshay Venkatesh</author>
<author pid="99/10550">Khaled Hamidouche</author>
<author pid="69/8165">Sreeram Potluri</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>Designing Optimized MPI Broadcast and Allreduce for Many Integrated Core (MIC) InfiniBand Clusters.</title>
<pages>63-70</pages>
<year>2013</year>
<booktitle>Hot Interconnects</booktitle>
<ee>https://doi.org/10.1109/HOTI.2013.26</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/HOTI.2013.26</ee>
<crossref>conf/hoti/2013</crossref>
<url>db/conf/hoti/hoti2013.html#KandallaVHPBP13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icpp/PotluriHVBP13" mdate="2023-03-24">
<author pid="69/8165">Sreeram Potluri</author>
<author pid="99/10550">Khaled Hamidouche</author>
<author pid="78/6299">Akshay Venkatesh</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>Efficient Inter-node MPI Communication Using GPUDirect RDMA for InfiniBand Clusters with NVIDIA GPUs.</title>
<pages>80-89</pages>
<year>2013</year>
<booktitle>ICPP</booktitle>
<ee>https://doi.org/10.1109/ICPP.2013.17</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICPP.2013.17</ee>
<crossref>conf/icpp/2013</crossref>
<url>db/conf/icpp/icpp2013.html#PotluriHVBP13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ipps/PotluriBWSP13" mdate="2023-03-24">
<author pid="69/8165">Sreeram Potluri</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="w/HaoWang-2">Hao Wang 0002</author>
<author pid="65/1958">Hari Subramoni</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>Extending OpenSHMEM for GPU Computing.</title>
<pages>1001-1012</pages>
<year>2013</year>
<booktitle>IPDPS</booktitle>
<ee>https://doi.org/10.1109/IPDPS.2013.104</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/IPDPS.2013.104</ee>
<crossref>conf/ipps/2013</crossref>
<url>db/conf/ipps/ipdps2013.html#PotluriBWSP13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/sc/PotluriBHVKSP13" mdate="2021-06-10">
<author pid="69/8165">Sreeram Potluri</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="99/10550">Khaled Hamidouche</author>
<author pid="78/6299">Akshay Venkatesh</author>
<author pid="91/7432">Krishna Chaitanya Kandalla</author>
<author pid="65/1958">Hari Subramoni</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>MVAPICH-PRISM: a proxy-based communication framework using InfiniBand and SCIF for intel MIC clusters.</title>
<pages>54:1-54:11</pages>
<year>2013</year>
<booktitle>SC</booktitle>
<ee>https://doi.org/10.1145/2503210.2503288</ee>
<crossref>conf/sc/2013</crossref>
<url>db/conf/sc/sc2013.html#PotluriBHVKSP13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ipps/PotluriWBSRP12" mdate="2023-03-24">
<author pid="69/8165">Sreeram Potluri</author>
<author pid="w/HaoWang-2">Hao Wang 0002</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="36/2597">Ashish Kumar Singh</author>
<author pid="56/3932">Carlos Rosales</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>Optimizing MPI Communication on Multi-GPU Systems Using CUDA Inter-Process Communication.</title>
<pages>1848-1857</pages>
<year>2012</year>
<booktitle>IPDPS Workshops</booktitle>
<ee>https://doi.org/10.1109/IPDPSW.2012.228</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/IPDPSW.2012.228</ee>
<crossref>conf/ipps/2012w</crossref>
<url>db/conf/ipps/ipdps2012w.html#PotluriWBSRP12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/pvm/BureddyWVPP12" mdate="2021-06-10">
<author pid="27/10106">Devendar Bureddy</author>
<author pid="w/HaoWang-2">Hao Wang 0002</author>
<author pid="78/6299">Akshay Venkatesh</author>
<author pid="69/8165">Sreeram Potluri</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>OMB-GPU: A Micro-Benchmark Suite for Evaluating MPI Libraries on GPU Clusters.</title>
<pages>110-120</pages>
<year>2012</year>
<booktitle>EuroMPI</booktitle>
<ee>https://doi.org/10.1007/978-3-642-33518-1_16</ee>
<crossref>conf/pvm/2012</crossref>
<url>db/conf/pvm/eurompi2012.html#BureddyWVPP12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/pvm/PotluriSBP11" mdate="2021-06-10">
<author pid="69/8165">Sreeram Potluri</author>
<author pid="41/591">Sayantan Sur</author>
<author pid="27/10106">Devendar Bureddy</author>
<author pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</author>
<title>Design and Implementation of Key Proposed MPI-3 One-Sided Communication Semantics on InfiniBand.</title>
<pages>321-324</pages>
<year>2011</year>
<booktitle>EuroMPI</booktitle>
<ee>https://doi.org/10.1007/978-3-642-24449-0_38</ee>
<crossref>conf/pvm/2011</crossref>
<url>db/conf/pvm/eurompi2011.html#PotluriSBP11</url>
</inproceedings>
</r>
<coauthors n="44" nc="1">
<co c="0"><na f="a/Aderholdt:Ferrol" pid="19/7905">Ferrol Aderholdt</na></co>
<co c="0"><na f="a/Arnold:Mark_Daniel" pid="131/4985">Mark Daniel Arnold</na></co>
<co c="0" n="2"><na f="b/Barth:William_L=" pid="08/1485">William L. Barth</na><na>Bill Barth</na></co>
<co c="0"><na f="b/Bloch:Gil" pid="99/2609">Gil Bloch</na></co>
<co c="0"><na f="c/Cho:David" pid="74/2156">David Cho</na></co>
<co c="0"><na f="d/Dubman:Mike" pid="129/4246">Mike Dubman</na></co>
<co c="0"><na f="e/Elias:George" pid="267/4443">George Elias</na></co>
<co c="0"><na f="g/Goldenberg:Dror" pid="07/1366">Dror Goldenberg</na></co>
<co c="0"><na f="g/Graham:Richard_L=" pid="86/5514">Richard L. Graham</na></co>
<co c="0"><na f="h/Hamidouche:Khaled" pid="99/10550">Khaled Hamidouche</na></co>
<co c="0"><na f="k/Kandalla:Krishna_Chaitanya" pid="91/7432">Krishna Chaitanya Kandalla</na></co>
<co c="0"><na f="k/Klein:Daniel" pid="93/6085">Daniel Klein</na></co>
<co c="0"><na f="k/Kotchubievsky:Sasha" pid="194/1376">Sasha Kotchubievsky</na></co>
<co c="0"><na f="k/Koushnir:Vladimir" pid="194/1355">Vladimir Koushnir</na></co>
<co c="0"><na f="l/Ladd:Joshua" pid="70/1942">Joshua Ladd</na></co>
<co c="0"><na f="l/Lebedev:Sergey" pid="25/9629">Sergey Lebedev</na></co>
<co c="0"><na f="l/Levi:Lion" pid="194/1364">Lion Levi</na></co>
<co c="0"><na f="l/Lui:Pak" pid="47/9777">Pak Lui</na></co>
<co c="0"><na f="m/Maor:Ophir" pid="267/4666">Ophir Maor</na></co>
<co c="0"><na f="m/Marelli:Ami" pid="267/4264">Ami Marelli</na></co>
<co c="0"><na f="m/Margolin:Alex" pid="194/1357">Alex Margolin</na></co>
<co c="0"><na f="p/Panda_0001:Dhabaleswar_K=" pid="p/DhabaleswarKPanda">Dhabaleswar K. Panda 0001</na></co>
<co c="0"><na f="p/Perkins:Jonathan_L=" pid="26/2675">Jonathan L. Perkins</na></co>
<co c="0"><na f="p/Petrov:Valentin" pid="76/300">Valentin Petrov</na></co>
<co c="0"><na f="p/Petrov:Valentine" pid="385/5435">Valentine Petrov</na></co>
<co c="0"><na f="p/Potluri:Sreeram" pid="69/8165">Sreeram Potluri</na></co>
<co c="0"><na f="q/Qin:Yong" pid="20/4298">Yong Qin</na></co>
<co c="0"><na f="r/Romlet:Evyatar" pid="267/4272">Evyatar Romlet</na></co>
<co c="0"><na f="r/Ronen:Tamir" pid="194/1384">Tamir Ronen</na></co>
<co c="0"><na f="r/Rosales:Carlos" pid="56/3932">Carlos Rosales</na></co>
<co c="0"><na f="r/Rosenstock:Hal" pid="194/1388">Hal Rosenstock</na></co>
<co c="0"><na f="s/Schulz:Karl_W=" pid="45/5668">Karl W. Schulz</na></co>
<co c="0"><na f="s/Shainer:Gilad" pid="52/6581">Gilad Shainer</na></co>
<co c="0"><na f="s/Shpiner:Alexander" pid="09/9457">Alexander Shpiner</na></co>
<co c="0"><na f="s/Singh:Ashish_Kumar" pid="36/2597">Ashish Kumar Singh</na></co>
<co c="0"><na f="s/Smith:Brian" pid="98/5275">Brian Smith</na></co>
<co c="0"><na f="s/Subramoni:Hari" pid="65/1958">Hari Subramoni</na></co>
<co c="0"><na f="s/Sur:Sayantan" pid="41/591">Sayantan Sur</na></co>
<co c="0"><na f="v/Venkata:Manjunath_Gorentla" pid="55/6076">Manjunath Gorentla Venkata</na></co>
<co c="0"><na f="v/Venkatesh:Akshay" pid="78/6299">Akshay Venkatesh</na></co>
<co c="0"><na f="w/Wang_0002:Hao" pid="w/HaoWang-2">Hao Wang 0002</na></co>
<co c="0"><na f="w/Wertheim:Oded" pid="194/1352">Oded Wertheim</na></co>
<co c="0"><na f="z/Zahavi:Eitan" pid="72/7892">Eitan Zahavi</na></co>
<co c="0"><na f="z/Zemah:Ido" pid="267/4288">Ido Zemah</na></co>
</coauthors>
</dblpperson>

