[{"@context":"https:\/\/schema.org\/","@type":"BlogPosting","@id":"https:\/\/blog.terabox.com\/insights\/eagle-eagle-2-aceleracion-inferencia-llm#BlogPosting","mainEntityOfPage":"https:\/\/blog.terabox.com\/insights\/eagle-eagle-2-aceleracion-inferencia-llm","headline":"EAGLE y EAGLE-2: Aceleraci\u00f3n de inferencia LLM sin p\u00e9rdida","name":"EAGLE y EAGLE-2: Aceleraci\u00f3n de inferencia LLM sin p\u00e9rdida","description":"\ud83d\udcfa V\u00eddeo de estudio recomendado hoy: https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg EAGLE: Aceleraci\u00f3n de LLM sin p\u00e9rdida mediante predicci\u00f3n de caracter\u00edsticas y \u00e1rboles din\u00e1micosEl problema de la inferencia secuencial y el muestreo especulativoEAGLE-1: Predicci\u00f3n de caracter\u00edsticas en lugar de tokensEAGLE-2: \u00c1rboles din\u00e1micos conscientes del... ","datePublished":"2026-08-20","dateModified":"2026-08-20","author":{"@type":"Person","@id":"https:\/\/blog.terabox.com\/author\/flextech-admin\/#Person","name":"flextech-admin","url":"https:\/\/blog.terabox.com\/author\/flextech-admin\/","image":{"@type":"ImageObject","@id":"https:\/\/secure.gravatar.com\/avatar\/ad516503a11cd5ca435acc9bb6523536?s=150&#038;d=mm&#038;r=gforcedefault=1","url":"https:\/\/secure.gravatar.com\/avatar\/ad516503a11cd5ca435acc9bb6523536?s=150&#038;d=mm&#038;r=gforcedefault=1","height":96,"width":96}},"publisher":{"@type":"Organization","name":"terabox","logo":{"@type":"ImageObject","@id":"http:\/\/blog.terabox.com\/wp-content\/uploads\/2021\/11\/logo\u4ea7\u54c1\u540d-\u7ad6\u7248.png","url":"http:\/\/blog.terabox.com\/wp-content\/uploads\/2021\/11\/logo\u4ea7\u54c1\u540d-\u7ad6\u7248.png","width":900,"height":900}},"image":{"@type":"ImageObject","@id":"https:\/\/blog.terabox.com\/wp-content\/uploads\/2026\/08\/mqdefault-14-1.jpg","url":"https:\/\/blog.terabox.com\/wp-content\/uploads\/2026\/08\/mqdefault-14-1.jpg","height":180,"width":320},"url":"https:\/\/blog.terabox.com\/insights\/eagle-eagle-2-aceleracion-inferencia-llm","video":{"@context":"http:\/\/schema.org\/","@type":"VideoObject","@id":"https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg#VideoObject","contentUrl":"https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg","name":"EAGLE and EAGLE-2: Lossless Inference Acceleration for LLMs - Hongyang Zhang","description":"About the seminar: https:\/\/faster-llms.vercel.app\n\nSpeaker: Hongyang Zhang (Waterloo & Vector Institute)\n\nTitle: EAGLE and EAGLE-2: Lossless Inference Acceleration for LLMs\n\nAbstract: This talk introduces the lossless large language model acceleration algorithm EAGLE and its follow-up, EAGLE-2 (\u201cEAGLE: Speculative Sampling Requires Rethinking Feature Uncertainty\u201d and \u201cEAGLE-2: Faster Inference of Language Models with Dynamic Draft Trees\u201d). EAGLE performs autoregression at a more structured feature level rather than at the token level, while incorporating sampling results to eliminate uncertainty. Thanks to these two innovations, EAGLE\u2019s draft model is both lightweight and accurate, improving the inference speed of large language models by 2.1x\u20133.8x while ensuring the output distribution remains unchanged, provably. EAGLE-2 introduces dynamic draft trees, leveraging the confidence of the draft model to approximate the acceptance rate of draft tokens and dynamically adjust the structure of the draft tree to increase the average acceptance length. Building on EAGLE-1, EAGLE-2 achieves an additional 20%-40% speed improvement, resulting in a total acceleration of 2.5x\u20135.0x while provably maintaining the original output distribution. EAGLE and EAGLE-2 have also been adopted in industry and open-sourced frameworks, integrated into platforms such as vLLM, SGLang, Intel LLM Library for PyTorch, Intel Extension for Transformers, and more.\n\nBio: Hongyang Zhang is a tenure-track assistant professor at University of Waterloo and Vector Institute for AI. He received his PhD in 2019 from the Machine Learning Department at Carnegie Mellon University and completed a Postdoc at Toyota Technological Institute at Chicago. He is the winner of NeurIPS 2018 Adversarial Vision Challenge, CVPR 2021 Security AI Challenger, AAAI New Faculty Highlights, Amazon Research Award, and WAIC Yunfan Award. He also regularly serves as an area chair for NeurIPS, ICLR, ICML, AISTATS, AAAI, ALT, and an action editor for DMLR.\n\nRecorded on Nov 25, 2024.","thumbnailUrl":["https:\/\/i.ytimg.com\/vi\/oXRSorx-Llg\/default.jpg","https:\/\/i.ytimg.com\/vi\/oXRSorx-Llg\/mqdefault.jpg","https:\/\/i.ytimg.com\/vi\/oXRSorx-Llg\/hqdefault.jpg","https:\/\/i.ytimg.com\/vi\/oXRSorx-Llg\/sddefault.jpg"],"uploadDate":"2025-02-01T18:16:46+00:00","duration":"PT48M26S","embedUrl":"https:\/\/www.youtube.com\/embed\/oXRSorx-Llg","publisher":{"@type":"Organization","@id":"https:\/\/www.youtube.com\/channel\/UCr0-QU5Jp82iHwSzTPQsJZA#Organization","url":"https:\/\/www.youtube.com\/channel\/UCr0-QU5Jp82iHwSzTPQsJZA","name":"Nadav Timor","description":"","logo":{"url":"https:\/\/yt3.ggpht.com\/Yc2eivbCJEMJhhbu-SyKnSSMrssk93PWgmZcfqAAF_BQ7sJT9FEgjGJr4CkJygg7V0Ytt7ol0Q=s800-c-k-c0x00ffffff-no-rj","width":800,"height":800,"@type":"ImageObject","@id":"https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg#VideoObject_publisher_logo_ImageObject"}},"potentialAction":{"@type":"SeekToAction","@id":"https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg#VideoObject_potentialAction","target":"https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg&t={seek_to_second_number}","startOffset-input":"required name=seek_to_second_number"},"interactionStatistic":[[{"@type":"InteractionCounter","@id":"https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg#VideoObject_interactionStatistic_WatchAction","interactionType":{"@type":"WatchAction"},"userInteractionCount":4490}],{"@type":"InteractionCounter","@id":"https:\/\/www.youtube.com\/watch?v=oXRSorx-Llg#VideoObject_interactionStatistic_LikeAction","interactionType":{"@type":"LikeAction"},"userInteractionCount":112}]},"about":["\u300eSpanish\u300f","Insights"],"wordCount":1868},{"@context":"https:\/\/schema.org\/","@type":"BreadcrumbList","itemListElement":[{"@type":"ListItem","position":1,"name":"Insights","item":"https:\/\/blog.terabox.com\/insights\/#breadcrumbitem"},{"@type":"ListItem","position":2,"name":"EAGLE y EAGLE-2: Aceleraci\u00f3n de inferencia LLM sin p\u00e9rdida","item":"https:\/\/blog.terabox.com\/insights\/eagle-eagle-2-aceleracion-inferencia-llm#breadcrumbitem"}]}]