CINXE.COM
🤗 Optimum
<!doctype html> <html class=""> <head> <meta charset="utf-8" /> <meta name="viewport" content="width=device-width, initial-scale=1.0, user-scalable=no" /> <meta name="description" content="We’re on a journey to advance and democratize artificial intelligence through open source and open science." /> <meta property="fb:app_id" content="1321688464574422" /> <meta name="twitter:card" content="summary_large_image" /> <meta name="twitter:site" content="@huggingface" /> <meta name="twitter:image" content="https://huggingface.co/front/thumbnails/docs/optimum.png" /> <meta property="og:title" content="🤗 Optimum" /> <meta property="og:type" content="website" /> <meta property="og:url" content="https://huggingface.co/docs/optimum/index" /> <meta property="og:image" content="https://huggingface.co/front/thumbnails/docs/optimum.png" /> <link rel="stylesheet" href="/front/build/kube-726083c/style.css" /> <link rel="preconnect" href="https://fonts.gstatic.com" /> <link href="https://fonts.googleapis.com/css2?family=Source+Sans+Pro:ital,wght@0,200;0,300;0,400;0,600;0,700;0,900;1,200;1,300;1,400;1,600;1,700;1,900&display=swap" rel="stylesheet" /> <link href="https://fonts.googleapis.com/css2?family=IBM+Plex+Mono:wght@400;600;700&display=swap" rel="stylesheet" /> <link rel="preload" href="https://cdnjs.cloudflare.com/ajax/libs/KaTeX/0.12.0/katex.min.css" as="style" onload="this.onload=null;this.rel='stylesheet'" /> <noscript> <link rel="stylesheet" href="https://cdnjs.cloudflare.com/ajax/libs/KaTeX/0.12.0/katex.min.css" /> </noscript> <link rel="canonical" href="https://huggingface.co/docs/optimum/index"> <link rel="alternate" hreflang="en" href="https://huggingface.co/docs/optimum/en/index"> <link rel="alternate" hreflang="x-default" href="https://huggingface.co/docs/optimum/index"> <title>🤗 Optimum</title> <script defer data-domain="huggingface.co" event-loggedIn="false" src="/js/script.pageview-props.js" ></script> <script> window.plausible = window.plausible || function () { (window.plausible.q = window.plausible.q || []).push(arguments); }; </script> <script> window.hubConfig = {"features":{"signupDisabled":false},"sshGitUrl":"git@hf.co","moonHttpUrl":"https:\/\/huggingface.co","captchaApiKey":"bd5f2066-93dc-4bdd-a64b-a24646ca3859","captchaDisabledOnSignup":false,"datasetViewerPublicUrl":"https:\/\/datasets-server.huggingface.co","stripePublicKey":"pk_live_x2tdjFXBCvXo2FFmMybezpeM00J6gPCAAc","environment":"production","userAgent":"HuggingFace (production)","spacesIframeDomain":"hf.space","spacesApiUrl":"https:\/\/api.hf.space","docSearchKey":"ece5e02e57300e17d152c08056145326e90c4bff3dd07d7d1ae40cf1c8d39cb6","logoDev":{"apiUrl":"https:\/\/img.logo.dev\/","apiKey":"pk_UHS2HZOeRnaSOdDp7jbd5w"}}; </script> <script type="text/javascript" src="https://de5282c3ca0c.edge.sdk.awswaf.com/de5282c3ca0c/526cf06acb0d/challenge.js" defer></script> </head> <body class="flex flex-col min-h-dvh bg-white dark:bg-gray-950 text-black DocBuilderPage"> <div class="flex min-h-dvh flex-col"> <div class="SVELTE_HYDRATER contents" data-target="MainHeader" data-props="{"classNames":"","isWide":true,"isZh":false}"><header class="border-b border-gray-100 "><div class="w-full px-4 flex h-16 items-center"><div class="flex flex-1 items-center"><a class="mr-5 flex flex-none items-center lg:mr-6" href="/"><img alt="Hugging Face's logo" class="w-7 md:mr-2" src="/front/assets/huggingface_logo-noborder.svg"> <span class="hidden whitespace-nowrap text-lg font-bold md:block">Hugging Face</span></a> <div class="relative flex-1 lg:max-w-sm mr-2 sm:mr-4 md:mr-3 xl:mr-6"><input autocomplete="off" class="w-full dark:bg-gray-950 pl-8 form-input-alt h-9 pr-3 focus:shadow-xl " name="" placeholder="Search models, datasets, users..." spellcheck="false" type="text" value=""> <svg class="absolute left-2.5 text-gray-400 top-1/2 transform -translate-y-1/2" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 32 32"><path d="M30 28.59L22.45 21A11 11 0 1 0 21 22.45L28.59 30zM5 14a9 9 0 1 1 9 9a9 9 0 0 1-9-9z" fill="currentColor"></path></svg> </div> <div class="flex flex-none items-center justify-center p-0.5 place-self-stretch lg:hidden"><button class="relative z-40 flex h-6 w-8 items-center justify-center" type="button"><svg width="1em" height="1em" viewBox="0 0 10 10" class="text-xl" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" preserveAspectRatio="xMidYMid meet" fill="currentColor"><path fill-rule="evenodd" clip-rule="evenodd" d="M1.65039 2.9999C1.65039 2.8066 1.80709 2.6499 2.00039 2.6499H8.00039C8.19369 2.6499 8.35039 2.8066 8.35039 2.9999C8.35039 3.1932 8.19369 3.3499 8.00039 3.3499H2.00039C1.80709 3.3499 1.65039 3.1932 1.65039 2.9999ZM1.65039 4.9999C1.65039 4.8066 1.80709 4.6499 2.00039 4.6499H8.00039C8.19369 4.6499 8.35039 4.8066 8.35039 4.9999C8.35039 5.1932 8.19369 5.3499 8.00039 5.3499H2.00039C1.80709 5.3499 1.65039 5.1932 1.65039 4.9999ZM2.00039 6.6499C1.80709 6.6499 1.65039 6.8066 1.65039 6.9999C1.65039 7.1932 1.80709 7.3499 2.00039 7.3499H8.00039C8.19369 7.3499 8.35039 7.1932 8.35039 6.9999C8.35039 6.8066 8.19369 6.6499 8.00039 6.6499H2.00039Z"></path></svg> </button> </div></div> <nav aria-label="Main" class="ml-auto hidden lg:block"><ul class="flex items-center space-x-1.5 2xl:space-x-2"><li class="hover:text-indigo-700"><a class="group flex items-center px-2 py-0.5 dark:hover:text-gray-400" href="/models"><svg class="mr-1.5 text-gray-400 group-hover:text-indigo-500" style="" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 24 24"><path class="uim-quaternary" d="M20.23 7.24L12 12L3.77 7.24a1.98 1.98 0 0 1 .7-.71L11 2.76c.62-.35 1.38-.35 2 0l6.53 3.77c.29.173.531.418.7.71z" opacity=".25" fill="currentColor"></path><path class="uim-tertiary" d="M12 12v9.5a2.09 2.09 0 0 1-.91-.21L4.5 17.48a2.003 2.003 0 0 1-1-1.73v-7.5a2.06 2.06 0 0 1 .27-1.01L12 12z" opacity=".5" fill="currentColor"></path><path class="uim-primary" d="M20.5 8.25v7.5a2.003 2.003 0 0 1-1 1.73l-6.62 3.82c-.275.13-.576.198-.88.2V12l8.23-4.76c.175.308.268.656.27 1.01z" fill="currentColor"></path></svg> Models</a> </li><li class="hover:text-red-700"><a class="group flex items-center px-2 py-0.5 dark:hover:text-gray-400" href="/datasets"><svg class="mr-1.5 text-gray-400 group-hover:text-red-500" style="" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 25 25"><ellipse cx="12.5" cy="5" fill="currentColor" fill-opacity="0.25" rx="7.5" ry="2"></ellipse><path d="M12.5 15C16.6421 15 20 14.1046 20 13V20C20 21.1046 16.6421 22 12.5 22C8.35786 22 5 21.1046 5 20V13C5 14.1046 8.35786 15 12.5 15Z" fill="currentColor" opacity="0.5"></path><path d="M12.5 7C16.6421 7 20 6.10457 20 5V11.5C20 12.6046 16.6421 13.5 12.5 13.5C8.35786 13.5 5 12.6046 5 11.5V5C5 6.10457 8.35786 7 12.5 7Z" fill="currentColor" opacity="0.5"></path><path d="M5.23628 12C5.08204 12.1598 5 12.8273 5 13C5 14.1046 8.35786 15 12.5 15C16.6421 15 20 14.1046 20 13C20 12.8273 19.918 12.1598 19.7637 12C18.9311 12.8626 15.9947 13.5 12.5 13.5C9.0053 13.5 6.06886 12.8626 5.23628 12Z" fill="currentColor"></path></svg> Datasets</a> </li><li class="hover:text-blue-700"><a class="group flex items-center px-2 py-0.5 dark:hover:text-gray-400" href="/spaces"><svg class="mr-1.5 text-gray-400 group-hover:text-blue-500" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" viewBox="0 0 25 25"><path opacity=".5" d="M6.016 14.674v4.31h4.31v-4.31h-4.31ZM14.674 14.674v4.31h4.31v-4.31h-4.31ZM6.016 6.016v4.31h4.31v-4.31h-4.31Z" fill="currentColor"></path><path opacity=".75" fill-rule="evenodd" clip-rule="evenodd" d="M3 4.914C3 3.857 3.857 3 4.914 3h6.514c.884 0 1.628.6 1.848 1.414a5.171 5.171 0 0 1 7.31 7.31c.815.22 1.414.964 1.414 1.848v6.514A1.914 1.914 0 0 1 20.086 22H4.914A1.914 1.914 0 0 1 3 20.086V4.914Zm3.016 1.102v4.31h4.31v-4.31h-4.31Zm0 12.968v-4.31h4.31v4.31h-4.31Zm8.658 0v-4.31h4.31v4.31h-4.31Zm0-10.813a2.155 2.155 0 1 1 4.31 0 2.155 2.155 0 0 1-4.31 0Z" fill="currentColor"></path><path opacity=".25" d="M16.829 6.016a2.155 2.155 0 1 0 0 4.31 2.155 2.155 0 0 0 0-4.31Z" fill="currentColor"></path></svg> Spaces</a> </li><li class="hover:text-yellow-700 max-xl:hidden"><a class="group flex items-center px-2 py-0.5 dark:hover:text-gray-400" href="/posts"><svg class="mr-1.5 text-gray-400 group-hover:text-yellow-500 !text-yellow-500" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" viewBox="0 0 12 12" preserveAspectRatio="xMidYMid meet"><path fill="currentColor" fill-rule="evenodd" d="M3.73 2.4A4.25 4.25 0 1 1 6 10.26H2.17l-.13-.02a.43.43 0 0 1-.3-.43l.01-.06a.43.43 0 0 1 .12-.22l.84-.84A4.26 4.26 0 0 1 3.73 2.4Z" clip-rule="evenodd"></path></svg> Posts</a> </li><li class="hover:text-yellow-700"><a class="group flex items-center px-2 py-0.5 dark:hover:text-gray-400" href="/docs"><svg xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" role="img" class="mr-1.5 text-gray-400 group-hover:text-yellow-500" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 32 32"><path opacity="0.5" d="M20.9022 5.10334L10.8012 10.8791L7.76318 9.11193C8.07741 8.56791 8.5256 8.11332 9.06512 7.7914L15.9336 3.73907C17.0868 3.08811 18.5002 3.26422 19.6534 3.91519L19.3859 3.73911C19.9253 4.06087 20.5879 4.56025 20.9022 5.10334Z" fill="currentColor"></path><path d="M10.7999 10.8792V28.5483C10.2136 28.5475 9.63494 28.4139 9.10745 28.1578C8.5429 27.8312 8.074 27.3621 7.74761 26.7975C7.42122 26.2327 7.24878 25.5923 7.24756 24.9402V10.9908C7.25062 10.3319 7.42358 9.68487 7.74973 9.1123L10.7999 10.8792Z" fill="currentColor" fill-opacity="0.75"></path><path fill-rule="evenodd" clip-rule="evenodd" d="M21.3368 10.8499V6.918C21.3331 6.25959 21.16 5.61234 20.8346 5.03949L10.7971 10.8727L10.8046 10.874L21.3368 10.8499Z" fill="currentColor"></path><path opacity="0.5" d="M21.7937 10.8488L10.7825 10.8741V28.5486L21.7937 28.5234C23.3344 28.5234 24.5835 27.2743 24.5835 25.7335V13.6387C24.5835 12.0979 23.4365 11.1233 21.7937 10.8488Z" fill="currentColor"></path></svg> Docs</a> </li><li class="hover:text-green-700"><a class="group flex items-center px-2 py-0.5 dark:hover:text-gray-400" href="/enterprise"><svg class="mr-1.5 text-gray-400 group-hover:text-green-500" xmlns="http://www.w3.org/2000/svg" fill="none" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 33 27"><path fill="currentColor" fill-rule="evenodd" d="M13.5.7a8.7 8.7 0 0 0-7.7 5.7L1 20.6c-1 3.1.9 5.7 4.1 5.7h15c3.3 0 6.8-2.6 7.8-5.7l4.6-14.2c1-3.1-.8-5.7-4-5.7h-15Zm1.1 5.7L9.8 20.3h9.8l1-3.1h-5.8l.8-2.5h4.8l1.1-3h-4.8l.8-2.3H23l1-3h-9.5Z" clip-rule="evenodd"></path></svg> Enterprise</a> </li> <li><a class="group flex items-center px-2 py-0.5 hover:text-gray-500 dark:hover:text-gray-400" href="/pricing">Pricing </a></li> <li><div class="relative group"> <button class="px-2 py-0.5 hover:text-gray-500 dark:hover:text-gray-600 flex items-center " type="button"> <svg class=" text-gray-500 w-5 group-hover:text-gray-400 dark:text-gray-300 dark:group-hover:text-gray-400" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" viewBox="0 0 32 18" preserveAspectRatio="xMidYMid meet"><path fill-rule="evenodd" clip-rule="evenodd" d="M14.4504 3.30221C14.4504 2.836 14.8284 2.45807 15.2946 2.45807H28.4933C28.9595 2.45807 29.3374 2.836 29.3374 3.30221C29.3374 3.76842 28.9595 4.14635 28.4933 4.14635H15.2946C14.8284 4.14635 14.4504 3.76842 14.4504 3.30221Z" fill="currentColor"></path><path fill-rule="evenodd" clip-rule="evenodd" d="M14.4504 9.00002C14.4504 8.53382 14.8284 8.15588 15.2946 8.15588H28.4933C28.9595 8.15588 29.3374 8.53382 29.3374 9.00002C29.3374 9.46623 28.9595 9.84417 28.4933 9.84417H15.2946C14.8284 9.84417 14.4504 9.46623 14.4504 9.00002Z" fill="currentColor"></path><path fill-rule="evenodd" clip-rule="evenodd" d="M14.4504 14.6978C14.4504 14.2316 14.8284 13.8537 15.2946 13.8537H28.4933C28.9595 13.8537 29.3374 14.2316 29.3374 14.6978C29.3374 15.164 28.9595 15.542 28.4933 15.542H15.2946C14.8284 15.542 14.4504 15.164 14.4504 14.6978Z" fill="currentColor"></path><path fill-rule="evenodd" clip-rule="evenodd" d="M1.94549 6.87377C2.27514 6.54411 2.80962 6.54411 3.13928 6.87377L6.23458 9.96907L9.32988 6.87377C9.65954 6.54411 10.194 6.54411 10.5237 6.87377C10.8533 7.20343 10.8533 7.73791 10.5237 8.06756L6.23458 12.3567L1.94549 8.06756C1.61583 7.73791 1.61583 7.20343 1.94549 6.87377Z" fill="currentColor"></path></svg> </button> </div></li> <li><hr class="h-5 w-0.5 border-none bg-gray-100 dark:bg-gray-800"></li> <li><a class="block cursor-pointer whitespace-nowrap px-2 py-0.5 hover:text-gray-500 dark:hover:text-gray-400" href="/login">Log In </a></li> <li><a class="whitespace-nowrap rounded-full border border-transparent bg-gray-900 px-3 py-1 leading-none text-white hover:border-black hover:bg-white hover:text-black" href="/join">Sign Up </a></li></ul></nav></div></header></div> <div class="SVELTE_HYDRATER contents" data-target="SSOBanner" data-props="{}"></div> <main class="flex flex-1 flex-col"><div class="relative lg:flex" id="hf-doc-container"><div class="sticky top-0 z-20 self-start"><div class="SVELTE_HYDRATER contents" data-target="SideMenu" data-props="{"chapters":[{"title":"Overview","isExpanded":true,"sections":[{"title":"🤗 Optimum","isExpanded":true,"id":"index","url":"/docs/optimum/index"},{"title":"Installation","isExpanded":true,"id":"installation","url":"/docs/optimum/installation"},{"title":"Quick tour","isExpanded":true,"id":"quicktour","url":"/docs/optimum/quicktour"},{"title":"Notebooks","isExpanded":true,"id":"notebooks","url":"/docs/optimum/notebooks"},{"title":"Conceptual guides","isExpanded":true,"sections":[{"title":"Quantization","isExpanded":true,"id":"concept_guides/quantization","url":"/docs/optimum/concept_guides/quantization"}]}]},{"title":"Nvidia","isExpanded":false,"sections":[{"title":"🤗 Optimum Nvidia","id":"nvidia_overview","url":"/docs/optimum/nvidia_overview"}]},{"title":"AMD","isExpanded":false,"sections":[{"title":"🤗 Optimum-AMD","id":"amd/index","url":"/docs/optimum/amd/index"},{"title":"Installation","id":"amd/installation","url":"/docs/optimum/amd/installation"},{"title":"AMD GPUs quicktour","id":"amd/amdgpu/overview","url":"/docs/optimum/amd/amdgpu/overview"},{"title":"Multi-GPU usage","id":"amd/amdgpu/perf_hardware","url":"/docs/optimum/amd/amdgpu/perf_hardware"},{"title":"Ryzen AI","isExpanded":false,"sections":[{"title":"Ryzen AI quicktour","id":"amd/ryzenai/overview","url":"/docs/optimum/amd/ryzenai/overview"},{"title":"How-to guides","isExpanded":false,"sections":[{"title":"How to apply quantization","id":"amd/ryzenai/usage_guides/quantization","url":"/docs/optimum/amd/ryzenai/usage_guides/quantization"},{"title":"Inference pipelines","id":"amd/ryzenai/usage_guides/pipelines","url":"/docs/optimum/amd/ryzenai/usage_guides/pipelines"}]},{"title":"Ryzen AI package reference","isExpanded":false,"sections":[{"title":"RyzenAI Models","id":"amd/ryzenai/package_reference/modeling","url":"/docs/optimum/amd/ryzenai/package_reference/modeling"},{"title":"RyzenAI Inference Pipelines","id":"amd/ryzenai/package_reference/pipelines","url":"/docs/optimum/amd/ryzenai/package_reference/pipelines"},{"title":"Configuration","id":"amd/ryzenai/package_reference/configuration","url":"/docs/optimum/amd/ryzenai/package_reference/configuration"},{"title":"Quantization","id":"amd/ryzenai/package_reference/quantization","url":"/docs/optimum/amd/ryzenai/package_reference/quantization"}]}]},{"title":"Brevitas","isExpanded":false,"sections":[{"title":"Usage guides","id":"amd/brevitas/usage_guide","url":"/docs/optimum/amd/brevitas/usage_guide"},{"title":"API reference","id":"amd/brevitas/api_reference","url":"/docs/optimum/amd/brevitas/api_reference"}]}]},{"title":"Intel","isExpanded":false,"sections":[{"title":"🤗 Optimum Intel","id":"intel/index","url":"/docs/optimum/intel/index"},{"title":"Installation","id":"intel/installation","url":"/docs/optimum/intel/installation"},{"title":"Neural Compressor","isExpanded":false,"sections":[{"title":"Optimization","id":"intel/neural_compressor/optimization","url":"/docs/optimum/intel/neural_compressor/optimization"},{"title":"Distributed Training","id":"intel/neural_compressor/distributed_training","url":"/docs/optimum/intel/neural_compressor/distributed_training"},{"title":"Reference","id":"intel/neural_compressor/reference","url":"/docs/optimum/intel/neural_compressor/reference"}]},{"title":"OpenVINO","isExpanded":false,"sections":[{"title":"Export","id":"intel/openvino/export","url":"/docs/optimum/intel/openvino/export"},{"title":"Inference","id":"intel/openvino/inference","url":"/docs/optimum/intel/openvino/inference"},{"title":"Optimization","id":"intel/openvino/optimization","url":"/docs/optimum/intel/openvino/optimization"},{"title":"Supported Models","id":"intel/openvino/models","url":"/docs/optimum/intel/openvino/models"},{"title":"Reference","id":"intel/openvino/reference","url":"/docs/optimum/intel/openvino/reference"},{"title":"Tutorials","isExpanded":false,"sections":[{"title":"Notebooks","id":"intel/openvino/tutorials/notebooks","url":"/docs/optimum/intel/openvino/tutorials/notebooks"},{"title":"Generate images with Diffusion models","id":"intel/openvino/tutorials/diffusers","url":"/docs/optimum/intel/openvino/tutorials/diffusers"}]}]},{"title":"IPEX","isExpanded":false,"sections":[{"title":"Inference","id":"intel/ipex/inference","url":"/docs/optimum/intel/ipex/inference"},{"title":"Supported Models","id":"intel/ipex/models","url":"/docs/optimum/intel/ipex/models"},{"title":"Tutorials","isExpanded":false,"sections":[{"title":"Notebooks","id":"intel/ipex/tutorials/notebooks","url":"/docs/optimum/intel/ipex/tutorials/notebooks"}]}]}]},{"title":"AWS Trainium/Inferentia","isExpanded":false,"sections":[{"title":"🤗 Optimum Neuron","id":"docs/optimum-neuron/index","url":"/docs/optimum/docs/optimum-neuron/index"}]},{"title":"Google TPUs","isExpanded":false,"sections":[{"title":"🤗 Optimum-TPU","id":"docs/optimum-tpu/index","url":"/docs/optimum/docs/optimum-tpu/index"}]},{"title":"for Intel Gaudi","isExpanded":false,"sections":[{"title":"🤗 Optimum for Intel Gaudi","id":"habana/index","url":"/docs/optimum/habana/index"},{"title":"Installation","id":"habana/installation","url":"/docs/optimum/habana/installation"},{"title":"Quickstart","id":"habana/quickstart","url":"/docs/optimum/habana/quickstart"},{"title":"Tutorials","isExpanded":false,"sections":[{"title":"Overview","id":"habana/tutorials/overview","url":"/docs/optimum/habana/tutorials/overview"},{"title":"Single-HPU Training","id":"habana/tutorials/single_hpu","url":"/docs/optimum/habana/tutorials/single_hpu"},{"title":"Distributed Training","id":"habana/tutorials/distributed","url":"/docs/optimum/habana/tutorials/distributed"},{"title":"Run Inference","id":"habana/tutorials/inference","url":"/docs/optimum/habana/tutorials/inference"},{"title":"Stable Diffusion","id":"habana/tutorials/stable_diffusion","url":"/docs/optimum/habana/tutorials/stable_diffusion"},{"title":"TGI on Gaudi","id":"habana/tutorials/tgi","url":"/docs/optimum/habana/tutorials/tgi"}]},{"title":"How-To Guides","isExpanded":false,"sections":[{"title":"Overview","id":"habana/usage_guides/overview","url":"/docs/optimum/habana/usage_guides/overview"},{"title":"Script Adaptation","id":"habana/usage_guides/script_adaptation","url":"/docs/optimum/habana/usage_guides/script_adaptation"},{"title":"Pretraining Transformers","id":"habana/usage_guides/pretraining","url":"/docs/optimum/habana/usage_guides/pretraining"},{"title":"Accelerating Training","id":"habana/usage_guides/accelerate_training","url":"/docs/optimum/habana/usage_guides/accelerate_training"},{"title":"Accelerating Inference","id":"habana/usage_guides/accelerate_inference","url":"/docs/optimum/habana/usage_guides/accelerate_inference"},{"title":"How to use DeepSpeed","id":"habana/usage_guides/deepspeed","url":"/docs/optimum/habana/usage_guides/deepspeed"},{"title":"Multi-node Training","id":"habana/usage_guides/multi_node_training","url":"/docs/optimum/habana/usage_guides/multi_node_training"},{"title":"Quantization","id":"habana/usage_guides/quantization","url":"/docs/optimum/habana/usage_guides/quantization"}]},{"title":"Reference","isExpanded":false,"sections":[{"title":"Gaudi Trainer","id":"habana/package_reference/trainer","url":"/docs/optimum/habana/package_reference/trainer"},{"title":"Gaudi Configuration","id":"habana/package_reference/gaudi_config","url":"/docs/optimum/habana/package_reference/gaudi_config"},{"title":"Distributed Runner","id":"habana/package_reference/distributed_runner","url":"/docs/optimum/habana/package_reference/distributed_runner"}]}]},{"title":"Furiosa","isExpanded":false,"sections":[{"title":"🤗 Optimum Furiosa","id":"furiosa/index","url":"/docs/optimum/furiosa/index"},{"title":"Installation","id":"furiosa/installation","url":"/docs/optimum/furiosa/installation"},{"title":"How-To Guides","isExpanded":false,"sections":[{"title":"Overview","id":"furiosa/usage_guides/overview","url":"/docs/optimum/furiosa/usage_guides/overview"},{"title":"Modeling","id":"furiosa/usage_guides/models","url":"/docs/optimum/furiosa/usage_guides/models"},{"title":"Quantization","id":"furiosa/usage_guides/quantization","url":"/docs/optimum/furiosa/usage_guides/quantization"}]},{"title":"Reference","isExpanded":false,"sections":[{"title":"Models","id":"furiosa/package_reference/modeling","url":"/docs/optimum/furiosa/package_reference/modeling"},{"title":"Configuration","id":"furiosa/package_reference/configuration","url":"/docs/optimum/furiosa/package_reference/configuration"},{"title":"Quantization","id":"furiosa/package_reference/quantization","url":"/docs/optimum/furiosa/package_reference/quantization"}]}]},{"title":"ONNX Runtime","isExpanded":false,"sections":[{"title":"Overview","id":"onnxruntime/overview","url":"/docs/optimum/onnxruntime/overview"},{"title":"Quick tour","id":"onnxruntime/quickstart","url":"/docs/optimum/onnxruntime/quickstart"},{"title":"How-to guides","isExpanded":false,"sections":[{"title":"Inference pipelines","id":"onnxruntime/usage_guides/pipelines","url":"/docs/optimum/onnxruntime/usage_guides/pipelines"},{"title":"Models for inference","id":"onnxruntime/usage_guides/models","url":"/docs/optimum/onnxruntime/usage_guides/models"},{"title":"How to apply graph optimization","id":"onnxruntime/usage_guides/optimization","url":"/docs/optimum/onnxruntime/usage_guides/optimization"},{"title":"How to apply dynamic and static quantization","id":"onnxruntime/usage_guides/quantization","url":"/docs/optimum/onnxruntime/usage_guides/quantization"},{"title":"How to accelerate training","id":"onnxruntime/usage_guides/trainer","url":"/docs/optimum/onnxruntime/usage_guides/trainer"},{"title":"Accelerated inference on NVIDIA GPUs","id":"onnxruntime/usage_guides/gpu","url":"/docs/optimum/onnxruntime/usage_guides/gpu"},{"title":"Accelerated inference on AMD GPUs","id":"onnxruntime/usage_guides/amdgpu","url":"/docs/optimum/onnxruntime/usage_guides/amdgpu"}]},{"title":"Conceptual guides","isExpanded":false,"sections":[{"title":"ONNX 🤝 ONNX Runtime","id":"onnxruntime/concept_guides/onnx","url":"/docs/optimum/onnxruntime/concept_guides/onnx"}]},{"title":"Reference","isExpanded":false,"sections":[{"title":"ONNX Runtime Models","id":"onnxruntime/package_reference/modeling_ort","url":"/docs/optimum/onnxruntime/package_reference/modeling_ort"},{"title":"Configuration","id":"onnxruntime/package_reference/configuration","url":"/docs/optimum/onnxruntime/package_reference/configuration"},{"title":"Optimization","id":"onnxruntime/package_reference/optimization","url":"/docs/optimum/onnxruntime/package_reference/optimization"},{"title":"Quantization","id":"onnxruntime/package_reference/quantization","url":"/docs/optimum/onnxruntime/package_reference/quantization"},{"title":"Trainer","id":"onnxruntime/package_reference/trainer","url":"/docs/optimum/onnxruntime/package_reference/trainer"}]}]},{"title":"Exporters","isExpanded":false,"sections":[{"title":"Overview","id":"exporters/overview","url":"/docs/optimum/exporters/overview"},{"title":"The TasksManager","id":"exporters/task_manager","url":"/docs/optimum/exporters/task_manager"},{"title":"ONNX","isExpanded":false,"sections":[{"title":"Overview","id":"exporters/onnx/overview","url":"/docs/optimum/exporters/onnx/overview"},{"title":"How-to guides","isExpanded":false,"sections":[{"title":"Export a model to ONNX","id":"exporters/onnx/usage_guides/export_a_model","url":"/docs/optimum/exporters/onnx/usage_guides/export_a_model"},{"title":"Add support for exporting an architecture to ONNX","id":"exporters/onnx/usage_guides/contribute","url":"/docs/optimum/exporters/onnx/usage_guides/contribute"}]},{"title":"Reference","isExpanded":false,"sections":[{"title":"ONNX configurations","id":"exporters/onnx/package_reference/configuration","url":"/docs/optimum/exporters/onnx/package_reference/configuration"},{"title":"Export functions","id":"exporters/onnx/package_reference/export","url":"/docs/optimum/exporters/onnx/package_reference/export"}]}]},{"title":"TFLite","isExpanded":false,"sections":[{"title":"Overview","id":"exporters/tflite/overview","url":"/docs/optimum/exporters/tflite/overview"},{"title":"How-to guides","isExpanded":false,"sections":[{"title":"Export a model to TFLite","id":"exporters/tflite/usage_guides/export_a_model","url":"/docs/optimum/exporters/tflite/usage_guides/export_a_model"},{"title":"Add support for exporting an architecture to TFLite","id":"exporters/tflite/usage_guides/contribute","url":"/docs/optimum/exporters/tflite/usage_guides/contribute"}]},{"title":"Reference","isExpanded":false,"sections":[{"title":"TFLite configurations","id":"exporters/tflite/package_reference/configuration","url":"/docs/optimum/exporters/tflite/package_reference/configuration"},{"title":"Export functions","id":"exporters/tflite/package_reference/export","url":"/docs/optimum/exporters/tflite/package_reference/export"}]}]}]},{"title":"BetterTransformer","isExpanded":false,"sections":[{"title":"Overview","id":"bettertransformer/overview","url":"/docs/optimum/bettertransformer/overview"},{"title":"Tutorials","isExpanded":false,"sections":[{"title":"Convert Transformers models to use BetterTransformer","id":"bettertransformer/tutorials/convert","url":"/docs/optimum/bettertransformer/tutorials/convert"},{"title":"How to add support for new architectures?","id":"bettertransformer/tutorials/contribute","url":"/docs/optimum/bettertransformer/tutorials/contribute"}]}]},{"title":"Torch FX","isExpanded":false,"sections":[{"title":"Overview","id":"torch_fx/overview","url":"/docs/optimum/torch_fx/overview"},{"title":"How-to guides","isExpanded":false,"sections":[{"title":"Optimization","id":"torch_fx/usage_guides/optimization","url":"/docs/optimum/torch_fx/usage_guides/optimization"}]},{"title":"Conceptual guides","isExpanded":false,"sections":[{"title":"Symbolic tracer","id":"torch_fx/concept_guides/symbolic_tracer","url":"/docs/optimum/torch_fx/concept_guides/symbolic_tracer"}]},{"title":"Reference","isExpanded":false,"sections":[{"title":"Optimization","id":"torch_fx/package_reference/optimization","url":"/docs/optimum/torch_fx/package_reference/optimization"}]}]},{"title":"LLM quantization","isExpanded":false,"sections":[{"title":"GPTQ quantization","id":"llm_quantization/usage_guides/quantization","url":"/docs/optimum/llm_quantization/usage_guides/quantization"}]},{"title":"Utilities","isExpanded":false,"sections":[{"title":"Dummy input generators","id":"utils/dummy_input_generators","url":"/docs/optimum/utils/dummy_input_generators"},{"title":"Normalized configurations","id":"utils/normalized_config","url":"/docs/optimum/utils/normalized_config"}]}],"chapterId":"index","docType":"docs","isLoggedIn":false,"lang":"en","langs":["en"],"library":"optimum","theme":"light","version":"main","versions":[{"version":"main"},{"version":"v1.23.3"},{"version":"v1.23.1"},{"version":"v1.22.0"},{"version":"v1.21.4"},{"version":"v1.21.2"},{"version":"v1.21.1"},{"version":"v1.21.0"},{"version":"v1.20.0"},{"version":"v1.19.0"},{"version":"v1.18.1"},{"version":"v1.18.0"},{"version":"v1.17.1"},{"version":"v1.16.2"},{"version":"v1.16.1"},{"version":"v1.16.0"},{"version":"v1.15.0"},{"version":"v1.14.0"},{"version":"v1.13.2"},{"version":"v1.13.1"},{"version":"v1.12.0"},{"version":"v1.11.2"},{"version":"v1.11.1"},{"version":"v1.10.1"},{"version":"v1.10.0"},{"version":"v1.9.0"},{"version":"v1.8.6"},{"version":"v1.8.5"},{"version":"v1.8.4"},{"version":"v1.8.3"},{"version":"v1.8.2"},{"version":"v1.8.1"},{"version":"v1.8.0"},{"version":"v1.7.3"},{"version":"v1.7.2"},{"version":"v1.7.1"},{"version":"v1.7.0"},{"version":"v1.6.4"},{"version":"v1.6.3"},{"version":"v1.6.2"},{"version":"v1.6.1"},{"version":"v1.6.0"},{"version":"v1.5.2"},{"version":"v1.5.1"},{"version":"v1.5.0"},{"version":"v1.4.1"},{"version":"v1.4.0"},{"version":"v1.3.0"},{"version":"v1.2.3"},{"version":"v1.2.2"},{"version":"v1.2.1"},{"version":"v1.2.0"},{"version":"v1.0.0"}],"title":"🤗 Optimum"}"> <div class="z-2 w-full flex-none lg:flex lg:h-dvh lg:w-[270px] lg:flex-col 2xl:w-[300px] false"><div class="shadow-alternate flex h-auto w-full items-center rounded-b-xl border-b bg-white py-2 text-lg leading-tight lg:hidden"> <div class="flex flex-1 cursor-pointer flex-col justify-center self-stretch pl-6"><p class="text-sm text-gray-400 first-letter:capitalize">Optimum documentation </p> <div class="mr-2 flex items-center"><p class="font-semibold">🤗 Optimum</p> <svg class="text-xl false" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 24 24"><path d="M16.293 9.293L12 13.586L7.707 9.293l-1.414 1.414L12 16.414l5.707-5.707z" fill="currentColor"></path></svg></div></div> <button class="hover:shadow-alternate group ml-auto mr-6 inline-flex flex-none cursor-pointer rounded-xl border p-2"><svg class="text-gray-500 group-hover:text-gray-700" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 32 32"><path d="M30 28.59L22.45 21A11 11 0 1 0 21 22.45L28.59 30zM5 14a9 9 0 1 1 9 9a9 9 0 0 1-9-9z" fill="currentColor"></path></svg></button></div> <div class="hidden flex-col justify-between border-b border-r bg-white bg-gradient-to-r p-4 lg:flex from-teal-50 to-white dark:from-gray-900 dark:to-gray-950"><div class="group relative mb-2 flex min-w-[50%] items-center self-start text-lg font-bold leading-tight first-letter:capitalize"><div class="mr-1.5 h-1.5 w-1.5 rounded-full bg-teal-500 flex-none"></div> <h1>Optimum</h1> <svg class="opacity-50 ml-0.5 flex-none group-hover:opacity-100" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 24 24"><path d="M16.293 9.293L12 13.586L7.707 9.293l-1.414 1.414L12 16.414l5.707-5.707z" fill="currentColor"></path></svg> <select class="absolute inset-0 border-none bg-white text-base opacity-0 outline-none"><option value="/docs">🏡 View all docs</option><option value="/docs/optimum-neuron" >AWS Trainium & Inferentia</option><option value="/docs/accelerate" >Accelerate</option><option value="/docs/sagemaker" >Amazon SageMaker</option><option value="https://argilla-io.github.io/argilla/" >Argilla</option><option value="/docs/autotrain" >AutoTrain</option><option value="/docs/bitsandbytes" >Bitsandbytes</option><option value="/docs/chat-ui" >Chat UI</option><option value="/docs/competitions" >Competitions</option><option value="/docs/dataset-viewer" >Dataset viewer</option><option value="/docs/datasets" >Datasets</option><option value="/docs/diffusers" >Diffusers</option><option value="https://distilabel.argilla.io/" >Distilabel</option><option value="/docs/evaluate" >Evaluate</option><option value="/docs/google-cloud" >Google Cloud</option><option value="/docs/optimum-tpu" >Google TPUs</option><option value="https://www.gradio.app/docs/" >Gradio</option><option value="/docs/hub" >Hub</option><option value="/docs/huggingface_hub" >Hub Python Library</option><option value="/docs/hugs" >Hugging Face Generative AI Services (HUGS)</option><option value="/docs/huggingface.js" >Huggingface.js</option><option value="/docs/api-inference" >Inference API (serverless)</option><option value="/docs/inference-endpoints" >Inference Endpoints (dedicated)</option><option value="/docs/leaderboards" >Leaderboards</option><option value="/docs/optimum" selected>Optimum</option><option value="/docs/peft" >PEFT</option><option value="/docs/safetensors" >Safetensors</option><option value="https://sbert.net/" >Sentence Transformers</option><option value="/docs/trl" >TRL</option><option value="/tasks" >Tasks</option><option value="/docs/text-embeddings-inference" >Text Embeddings Inference</option><option value="/docs/text-generation-inference" >Text Generation Inference</option><option value="/docs/tokenizers" >Tokenizers</option><option value="/docs/transformers" >Transformers</option><option value="/docs/transformers.js" >Transformers.js</option><option value="/docs/timm" >timm</option></select></div> <button class="shadow-alternate mb-2 flex w-full items-center rounded-full border bg-white px-2 py-1 text-left text-sm text-gray-400 ring-indigo-200 hover:bg-indigo-50 hover:ring-2 dark:border-gray-700 dark:ring-yellow-600 dark:hover:bg-gray-900 dark:hover:text-yellow-500"><svg class="flex-none mr-1.5" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 32 32"><path d="M30 28.59L22.45 21A11 11 0 1 0 21 22.45L28.59 30zM5 14a9 9 0 1 1 9 9a9 9 0 0 1-9-9z" fill="currentColor"></path></svg> <div>Search documentation</div> </button> <div class="flex items-center"> <select class="form-input !mt-0 mr-1 !w-20 rounded !border border-gray-200 p-1 text-xs uppercase dark:!text-gray-400"><option value="0" selected>main</option><option value="1" >v1.23.3</option><option value="2" >v1.22.0</option><option value="3" >v1.21.4</option><option value="4" >v1.20.0</option><option value="5" >v1.19.0</option><option value="6" >v1.18.1</option><option value="7" >v1.17.1</option><option value="8" >v1.16.2</option><option value="9" >v1.15.0</option><option value="10" >v1.14.0</option><option value="11" >v1.13.2</option><option value="12" >v1.12.0</option><option value="13" >v1.11.2</option><option value="14" >v1.10.1</option><option value="15" >v1.9.0</option><option value="16" >v1.8.6</option><option value="17" >v1.7.3</option><option value="18" >v1.6.4</option><option value="19" >v1.5.2</option><option value="20" >v1.4.1</option><option value="21" >v1.3.0</option><option value="22" >v1.2.3</option><option value="23" >v1.0.0</option></select> <select class="form-input mr-1 rounded border-gray-200 p-1 text-xs dark:!text-gray-400 !w-12 !mt-0 !border"><option value="en" selected>EN</option></select> <div class="relative inline-block"> <button class="rounded-full border border-gray-100 p-1.5 flex items-center text-sm text-gray-500 bg-white hover:bg-yellow-50 hover:border-yellow-200 dark:hover:bg-gray-800 dark:hover:border-gray-950 " type="button"> <svg class=" text-yellow-500" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 24 24" fill="currentColor"><path d="M6.05 4.14l-.39-.39a.993.993 0 0 0-1.4 0l-.01.01a.984.984 0 0 0 0 1.4l.39.39c.39.39 1.01.39 1.4 0l.01-.01a.984.984 0 0 0 0-1.4zM3.01 10.5H1.99c-.55 0-.99.44-.99.99v.01c0 .55.44.99.99.99H3c.56.01 1-.43 1-.98v-.01c0-.56-.44-1-.99-1zm9-9.95H12c-.56 0-1 .44-1 .99v.96c0 .55.44.99.99.99H12c.56.01 1-.43 1-.98v-.97c0-.55-.44-.99-.99-.99zm7.74 3.21c-.39-.39-1.02-.39-1.41-.01l-.39.39a.984.984 0 0 0 0 1.4l.01.01c.39.39 1.02.39 1.4 0l.39-.39a.984.984 0 0 0 0-1.4zm-1.81 15.1l.39.39a.996.996 0 1 0 1.41-1.41l-.39-.39a.993.993 0 0 0-1.4 0c-.4.4-.4 1.02-.01 1.41zM20 11.49v.01c0 .55.44.99.99.99H22c.55 0 .99-.44.99-.99v-.01c0-.55-.44-.99-.99-.99h-1.01c-.55 0-.99.44-.99.99zM12 5.5c-3.31 0-6 2.69-6 6s2.69 6 6 6s6-2.69 6-6s-2.69-6-6-6zm-.01 16.95H12c.55 0 .99-.44.99-.99v-.96c0-.55-.44-.99-.99-.99h-.01c-.55 0-.99.44-.99.99v.96c0 .55.44.99.99.99zm-7.74-3.21c.39.39 1.02.39 1.41 0l.39-.39a.993.993 0 0 0 0-1.4l-.01-.01a.996.996 0 0 0-1.41 0l-.39.39c-.38.4-.38 1.02.01 1.41z"></path></svg> </button> </div> <a href="https://github.com/huggingface/optimum" class="group ml-auto text-xs text-gray-500 hover:text-gray-700 hover:underline dark:hover:text-gray-300"><svg class="inline-block text-gray-500 group-hover:text-gray-700 dark:group-hover:text-gray-300 mr-1.5 -mt-1 w-4 h-4" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1.03em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 256 250"><path d="M128.001 0C57.317 0 0 57.307 0 128.001c0 56.554 36.676 104.535 87.535 121.46c6.397 1.185 8.746-2.777 8.746-6.158c0-3.052-.12-13.135-.174-23.83c-35.61 7.742-43.124-15.103-43.124-15.103c-5.823-14.795-14.213-18.73-14.213-18.73c-11.613-7.944.876-7.78.876-7.78c12.853.902 19.621 13.19 19.621 13.19c11.417 19.568 29.945 13.911 37.249 10.64c1.149-8.272 4.466-13.92 8.127-17.116c-28.431-3.236-58.318-14.212-58.318-63.258c0-13.975 5-25.394 13.188-34.358c-1.329-3.224-5.71-16.242 1.24-33.874c0 0 10.749-3.44 35.21 13.121c10.21-2.836 21.16-4.258 32.038-4.307c10.878.049 21.837 1.47 32.066 4.307c24.431-16.56 35.165-13.12 35.165-13.12c6.967 17.63 2.584 30.65 1.255 33.873c8.207 8.964 13.173 20.383 13.173 34.358c0 49.163-29.944 59.988-58.447 63.157c4.591 3.972 8.682 11.762 8.682 23.704c0 17.126-.148 30.91-.148 35.126c0 3.407 2.304 7.398 8.792 6.14C219.37 232.5 256 184.537 256 128.002C256 57.307 198.691 0 128.001 0zm-80.06 182.34c-.282.636-1.283.827-2.194.39c-.929-.417-1.45-1.284-1.15-1.922c.276-.655 1.279-.838 2.205-.399c.93.418 1.46 1.293 1.139 1.931zm6.296 5.618c-.61.566-1.804.303-2.614-.591c-.837-.892-.994-2.086-.375-2.66c.63-.566 1.787-.301 2.626.591c.838.903 1 2.088.363 2.66zm4.32 7.188c-.785.545-2.067.034-2.86-1.104c-.784-1.138-.784-2.503.017-3.05c.795-.547 2.058-.055 2.861 1.075c.782 1.157.782 2.522-.019 3.08zm7.304 8.325c-.701.774-2.196.566-3.29-.49c-1.119-1.032-1.43-2.496-.726-3.27c.71-.776 2.213-.558 3.315.49c1.11 1.03 1.45 2.505.701 3.27zm9.442 2.81c-.31 1.003-1.75 1.459-3.199 1.033c-1.448-.439-2.395-1.613-2.103-2.626c.301-1.01 1.747-1.484 3.207-1.028c1.446.436 2.396 1.602 2.095 2.622zm10.744 1.193c.036 1.055-1.193 1.93-2.715 1.95c-1.53.034-2.769-.82-2.786-1.86c0-1.065 1.202-1.932 2.733-1.958c1.522-.03 2.768.818 2.768 1.868zm10.555-.405c.182 1.03-.875 2.088-2.387 2.37c-1.485.271-2.861-.365-3.05-1.386c-.184-1.056.893-2.114 2.376-2.387c1.514-.263 2.868.356 3.061 1.403z" fill="currentColor"></path></svg> </a></div></div> <nav class="hidden flex-auto lg:flex bottom-0 left-0 w-full flex-col overflow-y-auto border-r px-4 pb-16 pt-3 text-[0.95rem] lg:w-[270px] 2xl:w-[300px]"> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] after:rotate-90 after:transform"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Overview<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="flex flex-col"><a class="rounded-xl bg-gradient-to-br from-black to-gray-900 py-1 pl-2 pr-2 text-white first:mt-1 last:mb-4 dark:from-gray-800 dark:to-gray-900 ml-2" href="/docs/optimum/index" id="index"><!-- HTML_TAG_START -->🤗 Optimum<!-- HTML_TAG_END --> </a><a class="transform py-1 pl-2 pr-2 text-gray-500 first:mt-1 last:mb-4 hover:translate-x-px hover:text-black dark:hover:text-gray-300 ml-2" href="/docs/optimum/installation" id="installation"><!-- HTML_TAG_START -->Installation<!-- HTML_TAG_END --> </a><a class="transform py-1 pl-2 pr-2 text-gray-500 first:mt-1 last:mb-4 hover:translate-x-px hover:text-black dark:hover:text-gray-300 ml-2" href="/docs/optimum/quicktour" id="quicktour"><!-- HTML_TAG_START -->Quick tour<!-- HTML_TAG_END --> </a><a class="transform py-1 pl-2 pr-2 text-gray-500 first:mt-1 last:mb-4 hover:translate-x-px hover:text-black dark:hover:text-gray-300 ml-2" href="/docs/optimum/notebooks" id="notebooks"><!-- HTML_TAG_START -->Notebooks<!-- HTML_TAG_END --> </a> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-2"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] after:rotate-90 after:transform"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Conceptual guides<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="flex flex-col"><a class="transform py-1 pl-2 pr-2 text-gray-500 first:mt-1 last:mb-4 hover:translate-x-px hover:text-black dark:hover:text-gray-300 ml-4" href="/docs/optimum/concept_guides/quantization" id="concept_guides/quantization"><!-- HTML_TAG_START -->Quantization<!-- HTML_TAG_END --> </a> </div> </div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Nvidia<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->AMD<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Intel<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->AWS Trainium/Inferentia<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Google TPUs<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->for Intel Gaudi<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Furiosa<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->ONNX Runtime<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Exporters<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->BetterTransformer<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Torch FX<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->LLM quantization<!-- HTML_TAG_END --></span> </span></span> </div></div> <div class="group flex cursor-pointer items-center pl-2 text-[0.8rem] font-semibold uppercase leading-9 hover:text-gray-700 dark:hover:text-gray-300 ml-0"><div class="flex after:absolute after:right-4 after:text-gray-500 group-hover:after:content-['▶'] false"><span><span class="inline-block space-x-1 leading-5"><span><!-- HTML_TAG_START -->Utilities<!-- HTML_TAG_END --></span> </span></span> </div></div> </nav></div></div></div> <div class="z-1 min-w-0 flex-1"><div class="flex justify-center bg-blue-50 px-6 py-1.5 text-xs text-blue-600 dark:bg-blue-900 dark:text-blue-200 sm:text-sm"><div>You are viewing <span class="italic">main</span> version, which requires <!-- HTML_TAG_START --><a class="underline" href=/docs/optimum/installation>installation from source</a><!-- HTML_TAG_END -->. If you'd like regular pip install, checkout the latest stable version (<a class="underline" href="/docs/optimum/v1.23.3/index">v1.23.3</a>). </div></div> <div class="px-6 pt-6 md:px-12 md:pb-16 md:pt-16"><div class="max-w-4xl mx-auto mb-10"><div class="relative overflow-hidden rounded-xl bg-gradient-to-br from-orange-300/10 px-4 py-5 ring-1 ring-orange-100/70 md:px-6 md:py-8"><img alt="Hugging Face's logo" class="absolute -bottom-6 -right-6 w-28 -rotate-45 md:hidden" src="/front/assets/huggingface_logo-noborder.svg"> <div class="mb-2 text-2xl font-bold dark:text-gray-200 md:mb-0">Join the Hugging Face community</div> <p class="mb-4 text-lg text-gray-400 dark:text-gray-300 md:mb-8">and get access to the augmented documentation experience </p> <div class="mb-8 hidden space-y-4 md:block xl:flex xl:space-x-6 xl:space-y-0"><div class="flex items-center"><div class="mr-3 flex h-9 w-9 flex-none items-center justify-center rounded-lg bg-gradient-to-br from-indigo-100 to-indigo-100/20 dark:to-indigo-100"><svg class="text-indigo-400 group-hover:text-indigo-500" style="" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 24 24"><path class="uim-quaternary" d="M20.23 7.24L12 12L3.77 7.24a1.98 1.98 0 0 1 .7-.71L11 2.76c.62-.35 1.38-.35 2 0l6.53 3.77c.29.173.531.418.7.71z" opacity=".25" fill="currentColor"></path><path class="uim-tertiary" d="M12 12v9.5a2.09 2.09 0 0 1-.91-.21L4.5 17.48a2.003 2.003 0 0 1-1-1.73v-7.5a2.06 2.06 0 0 1 .27-1.01L12 12z" opacity=".5" fill="currentColor"></path><path class="uim-primary" d="M20.5 8.25v7.5a2.003 2.003 0 0 1-1 1.73l-6.62 3.82c-.275.13-.576.198-.88.2V12l8.23-4.76c.175.308.268.656.27 1.01z" fill="currentColor"></path></svg></div> <div class="text-smd leading-tight text-gray-500 dark:text-gray-300 xl:max-w-[200px] 2xl:text-base">Collaborate on models, datasets and Spaces </div></div> <div class="flex items-center"><div class="mr-3 flex h-9 w-9 flex-none items-center justify-center rounded-lg bg-gradient-to-br from-orange-100 to-orange-100/20 dark:to-orange-50"><svg xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" class="text-xl text-yellow-400" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 24 24"><path d="M11 15H6l7-14v8h5l-7 14v-8z" fill="currentColor"></path></svg></div> <div class="text-smd leading-tight text-gray-500 dark:text-gray-300 xl:max-w-[200px] 2xl:text-base">Faster examples with accelerated inference </div></div> <div class="flex items-center"><div class="mr-3 flex h-9 w-9 flex-none items-center justify-center rounded-lg bg-gradient-to-br from-gray-500/10 to-gray-500/5"><svg class="text-gray-400" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" focusable="false" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 32 32"><path d="M14.9804 3C14.9217 3.0002 14.8631 3.00555 14.8054 3.016C11.622 3.58252 8.76073 5.30669 6.77248 7.85653C4.78422 10.4064 3.80955 13.6016 4.03612 16.8271C4.26268 20.0525 5.67447 23.0801 7.99967 25.327C10.3249 27.5738 13.3991 28.8811 16.6304 28.997C16.7944 29.003 16.9584 28.997 17.1204 28.997C19.2193 28.9984 21.2877 28.4943 23.1507 27.5274C25.0137 26.5605 26.6164 25.1592 27.8234 23.442C27.9212 23.294 27.9783 23.1229 27.9889 22.9458C27.9995 22.7687 27.9633 22.592 27.884 22.4333C27.8046 22.2747 27.6848 22.1397 27.5367 22.0421C27.3887 21.9444 27.2175 21.8875 27.0404 21.877C25.0426 21.7017 23.112 21.0693 21.3976 20.0288C19.6832 18.9884 18.231 17.5676 17.1533 15.8764C16.0756 14.1852 15.4011 12.2688 15.1822 10.2754C14.9632 8.28193 15.2055 6.26484 15.8904 4.38C15.9486 4.22913 15.97 4.06652 15.9527 3.90572C15.9354 3.74492 15.8799 3.59059 15.7909 3.45557C15.7019 3.32055 15.5819 3.20877 15.4409 3.12952C15.2999 3.05028 15.142 3.00587 14.9804 3Z" fill="currentColor"></path></svg></div> <div class="text-smd leading-tight text-gray-500 dark:text-gray-300 xl:max-w-[200px] 2xl:text-base">Switch between documentation themes </div></div></div> <div class="flex items-center space-x-2.5"><a href="/join"><button class="rounded-lg bg-white bg-gradient-to-br from-gray-100/20 to-gray-200/60 px-5 py-1.5 font-semibold text-gray-700 shadow-sm ring-1 ring-gray-300/60 hover:to-gray-100/70 hover:ring-gray-300/30 active:shadow-inner">Sign Up</button></a> <p class="text-gray-500 dark:text-gray-300">to get started</p></div></div></div> <div class="prose-doc prose relative mx-auto max-w-4xl break-words"><!-- HTML_TAG_START --> <link href="/docs/optimum/main/en/_app/immutable/assets/0.e3b0c442.css" rel="modulepreload"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/entry/start.9d426432.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/chunks/scheduler.6062bdaf.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/chunks/singletons.e8dbfdfe.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/chunks/paths.f8934bb5.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/entry/app.5d1e5a8d.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/chunks/index.4bca734e.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/nodes/0.85e9cc24.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/chunks/each.e59479a4.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/nodes/18.c3c30edb.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/chunks/Tip.b9ac1f03.js"> <link rel="modulepreload" href="/docs/optimum/main/en/_app/immutable/chunks/EditOnGithub.74ab2baa.js"><!-- HEAD_svelte-u9bgzb_START --><meta name="hf:doc:metadata" content="{"title":"🤗 Optimum","local":"-optimum","sections":[{"title":"Hardware partners","local":"hardware-partners","sections":[],"depth":2},{"title":"Open-source integrations","local":"open-source-integrations","sections":[],"depth":2}],"depth":1}"><!-- HEAD_svelte-u9bgzb_END --> <p></p> <h1 class="relative group"><a id="-optimum" class="header-link block pr-1.5 text-lg no-hover:hidden with-hover:absolute with-hover:p-1.5 with-hover:opacity-0 with-hover:group-hover:opacity-100 with-hover:right-full" href="#-optimum"><span><svg class="" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 256 256"><path d="M167.594 88.393a8.001 8.001 0 0 1 0 11.314l-67.882 67.882a8 8 0 1 1-11.314-11.315l67.882-67.881a8.003 8.003 0 0 1 11.314 0zm-28.287 84.86l-28.284 28.284a40 40 0 0 1-56.567-56.567l28.284-28.284a8 8 0 0 0-11.315-11.315l-28.284 28.284a56 56 0 0 0 79.196 79.197l28.285-28.285a8 8 0 1 0-11.315-11.314zM212.852 43.14a56.002 56.002 0 0 0-79.196 0l-28.284 28.284a8 8 0 1 0 11.314 11.314l28.284-28.284a40 40 0 0 1 56.568 56.567l-28.285 28.285a8 8 0 0 0 11.315 11.314l28.284-28.284a56.065 56.065 0 0 0 0-79.196z" fill="currentColor"></path></svg></span></a> <span>🤗 Optimum</span></h1> <p data-svelte-h="svelte-2e1kxa">🤗 Optimum is an extension of <a href="https://huggingface.co/docs/transformers" rel="nofollow">Transformers</a> that provides a set of performance optimization tools to train and run models on targeted hardware with maximum efficiency.</p> <p data-svelte-h="svelte-1199uaa">The AI ecosystem evolves quickly, and more and more specialized hardware along with their own optimizations are emerging every day. As such, Optimum enables developers to efficiently use any of these platforms with the same ease inherent to Transformers.</p> <p data-svelte-h="svelte-u6or0s">🤗 Optimum is distributed as a collection of packages - check out the links below for an in-depth look at each one.</p> <h2 class="relative group"><a id="hardware-partners" class="header-link block pr-1.5 text-lg no-hover:hidden with-hover:absolute with-hover:p-1.5 with-hover:opacity-0 with-hover:group-hover:opacity-100 with-hover:right-full" href="#hardware-partners"><span><svg class="" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 256 256"><path d="M167.594 88.393a8.001 8.001 0 0 1 0 11.314l-67.882 67.882a8 8 0 1 1-11.314-11.315l67.882-67.881a8.003 8.003 0 0 1 11.314 0zm-28.287 84.86l-28.284 28.284a40 40 0 0 1-56.567-56.567l28.284-28.284a8 8 0 0 0-11.315-11.315l-28.284 28.284a56 56 0 0 0 79.196 79.197l28.285-28.285a8 8 0 1 0-11.315-11.314zM212.852 43.14a56.002 56.002 0 0 0-79.196 0l-28.284 28.284a8 8 0 1 0 11.314 11.314l28.284-28.284a40 40 0 0 1 56.568 56.567l-28.285 28.285a8 8 0 0 0 11.315 11.314l28.284-28.284a56.065 56.065 0 0 0 0-79.196z" fill="currentColor"></path></svg></span></a> <span>Hardware partners</span></h2> <p data-svelte-h="svelte-11fz6d7">The packages below enable you to get the best of the 🤗 Hugging Face ecosystem on various types of devices.</p> <div class="mt-10" data-svelte-h="svelte-1ngloan"><div class="w-full flex flex-col space-y-4 md:space-y-0 md:grid md:grid-cols-4 md:gap-y-4 md:gap-x-5"><a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="https://github.com/huggingface/optimum-nvidia"><div class="w-full text-center bg-gradient-to-br from-green-600 to-green-600 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">NVIDIA</div> <p class="text-gray-700">Accelerate inference with NVIDIA TensorRT-LLM on the <span class="underline" onclick="event.preventDefault(); window.open('https://developer.nvidia.com/blog/nvidia-tensorrt-llm-supercharges-large-language-model-inference-on-nvidia-h100-gpus/', '_blank');">NVIDIA platform</span></p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./amd/index"><div class="w-full text-center bg-gradient-to-br from-red-600 to-red-600 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">AMD</div> <p class="text-gray-700">Enable performance optimizations for <span class="underline" onclick="event.preventDefault(); window.open('https://www.amd.com/en/graphics/instinct-server-accelerators', '_blank');">AMD Instinct GPUs</span> and <span class="underline" onclick="event.preventDefault(); window.open('https://ryzenai.docs.amd.com/en/latest/index.html', '_blank');">AMD Ryzen AI NPUs</span></p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./intel/index"><div class="w-full text-center bg-gradient-to-br from-blue-400 to-blue-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">Intel</div> <p class="text-gray-700">Optimize your model to speedup inference with <span class="underline" onclick="event.preventDefault(); window.open('https://docs.openvino.ai/latest/index.html', '_blank');">OpenVINO</span> , <span class="underline" onclick="event.preventDefault(); window.open('https://www.intel.com/content/www/us/en/developer/tools/oneapi/neural-compressor.html', '_blank');">Neural Compressor</span> and <span class="underline" onclick="event.preventDefault(); window.open('https://intel.github.io/intel-extension-for-pytorch/index.html', '_blank');">IPEX</span></p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="https://huggingface.co/docs/optimum-neuron/index"><div class="w-full text-center bg-gradient-to-br from-orange-400 to-orange-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">AWS Trainium/Inferentia</div> <p class="text-gray-700">Accelerate your training and inference workflows with <span class="underline" onclick="event.preventDefault(); window.open('https://aws.amazon.com/machine-learning/trainium/', '_blank');">AWS Trainium</span> and <span class="underline" onclick="event.preventDefault(); window.open('https://aws.amazon.com/machine-learning/inferentia/', '_blank');">AWS Inferentia</span></p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="https://huggingface.co/docs/optimum-tpu/index"><div class="w-full text-center bg-gradient-to-tr from-blue-200 to-blue-600 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">Google TPUs</div> <p class="text-gray-700">Accelerate your training and inference workflows with <span class="underline" onclick="event.preventDefault(); window.open('https://cloud.google.com/tpu', '_blank');">Google TPUs</span></p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./habana/index"><div class="w-full text-center bg-gradient-to-br from-indigo-400 to-indigo-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">Habana</div> <p class="text-gray-700">Maximize training throughput and efficiency with <span class="underline" onclick="event.preventDefault(); window.open('https://docs.habana.ai/en/latest/Gaudi_Overview/Gaudi_Architecture.html', '_blank');">Habana's Gaudi processor</span></p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./furiosa/index"><div class="w-full text-center bg-gradient-to-br from-green-400 to-green-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">FuriosaAI</div> <p class="text-gray-700">Fast and efficient inference on <span class="underline" onclick="event.preventDefault(); window.open('https://www.furiosa.ai/', '_blank');">FuriosaAI WARBOY</span></p></a></div></div> <div class="course-tip bg-gradient-to-br dark:bg-gradient-to-r before:border-green-500 dark:before:border-green-800 from-green-50 dark:from-gray-900 to-white dark:to-gray-950 border border-green-50 text-green-700 dark:text-gray-400"><p data-svelte-h="svelte-hc7a16">Some packages provide hardware-agnostic features (e.g. INC interface in Optimum Intel).</p></div> <h2 class="relative group"><a id="open-source-integrations" class="header-link block pr-1.5 text-lg no-hover:hidden with-hover:absolute with-hover:p-1.5 with-hover:opacity-0 with-hover:group-hover:opacity-100 with-hover:right-full" href="#open-source-integrations"><span><svg class="" xmlns="http://www.w3.org/2000/svg" xmlns:xlink="http://www.w3.org/1999/xlink" aria-hidden="true" role="img" width="1em" height="1em" preserveAspectRatio="xMidYMid meet" viewBox="0 0 256 256"><path d="M167.594 88.393a8.001 8.001 0 0 1 0 11.314l-67.882 67.882a8 8 0 1 1-11.314-11.315l67.882-67.881a8.003 8.003 0 0 1 11.314 0zm-28.287 84.86l-28.284 28.284a40 40 0 0 1-56.567-56.567l28.284-28.284a8 8 0 0 0-11.315-11.315l-28.284 28.284a56 56 0 0 0 79.196 79.197l28.285-28.285a8 8 0 1 0-11.315-11.314zM212.852 43.14a56.002 56.002 0 0 0-79.196 0l-28.284 28.284a8 8 0 1 0 11.314 11.314l28.284-28.284a40 40 0 0 1 56.568 56.567l-28.285 28.285a8 8 0 0 0 11.315 11.314l28.284-28.284a56.065 56.065 0 0 0 0-79.196z" fill="currentColor"></path></svg></span></a> <span>Open-source integrations</span></h2> <p data-svelte-h="svelte-kk9o8p">🤗 Optimum also supports a variety of open-source frameworks to make model optimization very easy.</p> <div class="mt-10" data-svelte-h="svelte-1usc1uy"><div class="w-full flex flex-col space-y-4 md:space-y-0 md:grid md:grid-cols-3 md:gap-y-4 md:gap-x-5"><a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./onnxruntime/overview"><div class="w-full text-center bg-gradient-to-br from-pink-400 to-pink-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">ONNX Runtime</div> <p class="text-gray-700">Apply quantization and graph optimization to accelerate Transformers models training and inference with <span class="underline" onclick="event.preventDefault(); window.open('https://onnxruntime.ai/', '_blank');">ONNX Runtime</span></p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./exporters/overview"><div class="w-full text-center bg-gradient-to-br from-purple-400 to-purple-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">Exporters</div> <p class="text-gray-700">Export your PyTorch or TensorFlow model to different formats such as ONNX and TFLite</p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./bettertransformer/overview"><div class="w-full text-center bg-gradient-to-br from-yellow-400 to-yellow-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">BetterTransformer</div> <p class="text-gray-700">A one-liner integration to use <span class="underline" onclick="event.preventDefault(); window.open('https://pytorch.org/blog/a-better-transformer-for-fast-transformer-encoder-inference/', '_blank');">PyTorch's BetterTransformer</span> with Transformers models</p></a> <a class="!no-underline border dark:border-gray-700 p-5 rounded-lg shadow hover:shadow-lg" href="./torch_fx/overview"><div class="w-full text-center bg-gradient-to-br from-green-400 to-green-500 rounded-lg py-1.5 font-semibold mb-5 text-white text-lg leading-relaxed">Torch FX</div> <p class="text-gray-700">Create and compose custom graph transformations to optimize PyTorch Transformers models with <span class="underline" onclick="event.preventDefault(); window.open('https://pytorch.org/docs/stable/fx.html#', '_blank');">Torch FX</span></p></a></div></div> <a class="!text-gray-400 !no-underline text-sm flex items-center not-prose mt-4" href="https://github.com/huggingface/optimum/blob/main/docs/source/index.mdx" target="_blank"><span data-svelte-h="svelte-1kd6by1"><</span> <span data-svelte-h="svelte-x0xyl0">></span> <span data-svelte-h="svelte-1dajgef"><span class="underline ml-1.5">Update</span> on GitHub</span></a> <p></p> <script> { __sveltekit_9h6edu = { assets: "/docs/optimum/main/en", base: "/docs/optimum/main/en", env: {} }; const element = document.currentScript.parentElement; const data = [null,null]; Promise.all([ import("/docs/optimum/main/en/_app/immutable/entry/start.9d426432.js"), import("/docs/optimum/main/en/_app/immutable/entry/app.5d1e5a8d.js") ]).then(([kit, app]) => { kit.start(app, element, { node_ids: [0, 18], data, form: null, error: null }); }); } </script> <!-- HTML_TAG_END --></div> <div class="SVELTE_HYDRATER contents" data-target="DocFooterNav" data-props="{"classNames":"mx-auto mt-16 flex max-w-4xl items-center pb-8 font-sans font-medium leading-6 xl:mt-32","chapterNext":{"title":"Installation","isExpanded":true,"id":"installation","url":"/docs/optimum/installation"},"isCourse":false,"isLoggedIn":false}"><div class="mx-auto mt-16 flex max-w-4xl items-center pb-8 font-sans font-medium leading-6 xl:mt-32"> <a href="/docs/optimum/installation" class="ml-auto flex transform items-center text-right text-gray-600 transition-all hover:translate-x-px hover:text-gray-900 dark:hover:text-gray-300">Installation<span class="ml-2 translate-y-px">→</span></a></div></div></div></div> <div class="sticky top-0 self-start"><div class="SVELTE_HYDRATER contents" data-target="SubSideMenu" data-props="{"chapter":{"title":"🤗 Optimum","isExpanded":true,"id":"-optimum","url":"#-optimum","sections":[{"title":"Hardware partners","isExpanded":true,"id":"hardware-partners","url":"#hardware-partners","sections":[]},{"title":"Open-source integrations","isExpanded":true,"id":"open-source-integrations","url":"#open-source-integrations","sections":[]}]}}"> <nav class="hidden h-dvh w-[270px] flex-none flex-col space-y-3 overflow-y-auto break-words border-l pb-16 pl-6 pr-10 pt-24 text-sm lg:flex 2xl:w-[305px]"> <a href="#-optimum" class=" text-gray-400 transform hover:translate-x-px hover:text-gray-700 dark:hover:text-gray-300" id="nav--optimum"><!-- HTML_TAG_START -->🤗 <wbr>Optimum<!-- HTML_TAG_END --></a> <a href="#hardware-partners" class="pl-4 text-gray-400 transform hover:translate-x-px hover:text-gray-700 dark:hover:text-gray-300" id="nav-hardware-partners"><!-- HTML_TAG_START --><wbr>Hardware partners<!-- HTML_TAG_END --></a> <a href="#open-source-integrations" class="pl-4 text-gray-400 transform hover:translate-x-px hover:text-gray-700 dark:hover:text-gray-300" id="nav-open-source-integrations"><!-- HTML_TAG_START --><wbr>Open-source integrations<!-- HTML_TAG_END --></a> </nav></div></div></div> <div id="doc-footer"></div></main> </div> <script> import("\/front\/build\/kube-726083c\/index.js"); window.moonSha = "kube-726083c\/"; window.__hf_deferred = {}; </script> <!-- Stripe --> <script> if (["hf.co", "huggingface.co"].includes(window.location.hostname)) { const script = document.createElement("script"); script.src = "https://js.stripe.com/v3/"; script.async = true; document.head.appendChild(script); } </script> </body> </html>