<!DOCTYPE html><html lang="en" class="lenis lenis-smooth" data-astro-cid-sckkx6r4> <head><meta charset="UTF-8"><meta name="viewport" content="width=device-width, initial-scale=1.0"><meta name="theme-color" content="#0a0a0a"><!-- SEO --><title>Startups | Launch with confidence | Future AGI</title><meta name="description" content="$6K in free credits, 6 months Pro access, and direct engineering support. Build AI products that don't hallucinate."><link rel="canonical" href="https://futureagi.com/startups/"><meta name="robots" content="index, follow"><meta name="msvalidate.01" content="B8799A595FEACE500CE90BE77714035D"><!-- Open Graph --><meta property="og:type" content="website"><meta property="og:url" content="https://futureagi.com/startups/"><meta property="og:title" content="Startups | Launch with confidence | Future AGI"><meta property="og:description" content="$6K in free credits, 6 months Pro access, and direct engineering support. Build AI products that don't hallucinate."><meta property="og:site_name" content="Future AGI"><meta property="og:locale" content="en_US"><meta property="og:image" content="https://futureagi.com/og-image.png"><meta property="og:image:width" content="1200"><meta property="og:image:height" content="630"><meta property="og:image:type" content="image/png"><!-- Twitter / X --><meta name="twitter:card" content="summary_large_image"><meta name="twitter:site" content="@futureagi"><meta name="twitter:title" content="Startups | Launch with confidence | Future AGI"><meta name="twitter:description" content="$6K in free credits, 6 months Pro access, and direct engineering support. Build AI products that don't hallucinate."><meta name="twitter:image" content="https://futureagi.com/og-image.png"><!-- Favicon --><link rel="icon" type="image/svg+xml" href="/favicon.svg"><link rel="apple-touch-icon" sizes="180x180" href="/apple-touch-icon.png"><!-- RSS --><link rel="alternate" type="application/rss+xml" title="Future AGI Changelog" href="https://futureagi.com/changelog/rss.xml"><!-- Homepage structured data --><!-- Platform page structured data (auto-injected for /platform/*) --><!-- Preconnect: critical-path origins (≤4 to avoid wasting sockets) --><link rel="preconnect" href="https://fonts.googleapis.com"><link rel="preconnect" href="https://fonts.gstatic.com" crossorigin><link rel="preconnect" href="https://cdn.simpleicons.org"><!-- DNS prefetch: non-critical origins --><link rel="dns-prefetch" href="https://fonts.googleapis.com"><link rel="dns-prefetch" href="https://fonts.gstatic.com"><link rel="dns-prefetch" href="https://js.hs-scripts.com"><link rel="dns-prefetch" href="https://js.hs-analytics.net"><link rel="dns-prefetch" href="https://46657045.hs-sites.com"><link rel="dns-prefetch" href="https://zd.futureagi.com"><link rel="dns-prefetch" href="https://www.googletagmanager.com"><link rel="dns-prefetch" href="https://pixel-config.reddit.com"><link rel="dns-prefetch" href="https://alb.reddit.com"><link rel="dns-prefetch" href="https://tags.clickagy.com"><link rel="preload" as="style" href="https://fonts.googleapis.com/css2?family=Inter:wght@300;400;500;600;700&family=JetBrains+Mono:wght@400;500&display=swap"><link href="https://fonts.googleapis.com/css2?family=Inter:wght@300;400;500;600;700&family=JetBrains+Mono:wght@400;500&display=swap" rel="stylesheet" media="print" onload="this.media='all'"><noscript><link href="https://fonts.googleapis.com/css2?family=Inter:wght@300;400;500;600;700&family=JetBrains+Mono:wght@400;500&display=swap" rel="stylesheet"></noscript><!-- Structured Data: Organization + WebSite (on every page) --><script type="application/ld+json">{"@context":"https://schema.org","@type":"Organization","@id":"https://futureagi.com/#organization","name":"Future AGI","url":"https://futureagi.com","description":"The complete platform to test, guard, and monitor AI agents — hallucination detection, RAG observability, and AI agent evaluation.","sameAs":["https://x.com/FutureAGI_","https://linkedin.com/company/futureagi","https://github.com/future-agi","https://discord.com/invite/n2tCUKBkAw"],"logo":{"@type":"ImageObject","url":"https://futureagi.com/favicon.svg","width":112,"height":112},"contactPoint":{"@type":"ContactPoint","email":"hello@futureagi.com","contactType":"customer support"}}</script><script type="application/ld+json">{"@context":"https://schema.org","@type":"WebSite","@id":"https://futureagi.com/#website","url":"https://futureagi.com","name":"Future AGI","publisher":{"@id":"https://futureagi.com/#organization"},"potentialAction":{"@type":"SearchAction","target":{"@type":"EntryPoint","urlTemplate":"https://docs.futureagi.com/docs?q={search_term_string}"},"query-input":"required name=search_term_string"}}</script><!-- All ad-tracking (gtag, Reddit, Twitter, attribution capture, consent,
         adblock detection) lives in one shared partial used by every layout. --><!-- gtag Consent Mode v2 — must run BEFORE any gtag('config') so defaults
     apply correctly. Outside EEA/UK: granted. Inside EEA/UK: denied until
     user accepts via the banner below. Google's IP-based region matching
     handles this server-side; the banner gives EEA users a way to upgrade. --><script>
  window.dataLayer = window.dataLayer || [];
  function gtag(){dataLayer.push(arguments);}
  window.gtag = gtag;

  // Default: granted everywhere
  gtag('consent', 'default', {
    ad_storage: 'granted',
    ad_user_data: 'granted',
    ad_personalization: 'granted',
    analytics_storage: 'granted',
    wait_for_update: 500
  });
  // Override: denied in EEA/UK regions until user accepts
  gtag('consent', 'default', {
    ad_storage: 'denied',
    ad_user_data: 'denied',
    ad_personalization: 'denied',
    analytics_storage: 'denied',
    region: ['AT','BE','BG','HR','CY','CZ','DK','EE','FI','FR','DE','GR','HU','IE','IT','LV','LT','LU','MT','NL','PL','PT','RO','SK','SI','ES','SE','GB','IS','LI','NO','CH'],
    wait_for_update: 500
  });
</script> <!-- Google Analytics + Google Ads — stub gtag synchronously so calls before
     the real script lands get queued; gtag.js loads on requestIdleCallback
     so pixel-helper extensions detect it on first load without blocking LCP.
     Cross-domain linker carries gclid/client ID across subdomains. --><script>(function(){const GA_ID = "G-BFBK89532X";
const GOOGLE_ADS_ID = "AW-18081760643";
const LINKER_DOMAINS = ["futureagi.com","docs.futureagi.com","app.futureagi.com"];

  gtag('js', new Date());
  if (GOOGLE_ADS_ID) gtag('config', GOOGLE_ADS_ID, {
    linker: { domains: LINKER_DOMAINS, accept_incoming: true }
  });
  if (GA_ID) gtag('config', GA_ID, {
    send_page_view: true,
    linker: { domains: LINKER_DOMAINS, accept_incoming: true }
  });
  var _gtagLoaded = false;
  window._loadGtag = function () {
    if (_gtagLoaded || !GOOGLE_ADS_ID) return;
    _gtagLoaded = true;
    var s = document.createElement('script');
    s.async = true;
    s.src = 'https://www.googletagmanager.com/gtag/js?id=' + GOOGLE_ADS_ID;
    document.head.appendChild(s);
  };
  (window.requestIdleCallback || function(cb){ return setTimeout(cb, 200); })(window._loadGtag, { timeout: 2000 });
})();</script> <!-- Capture ad-click attribution (gclid, utm_*) into a .futureagi.com
     cookie so the app subdomain reads it on signup. Runs on every page;
     only an actual ad click (gclid/gbraid/wbraid) overwrites prior. --><script type="module" src="/_astro/AdPixels.astro_astro_type_script_index_0_lang.Baj_DXh8.js"></script> <!-- Reddit Pixel — idle-time load so Reddit Pixel Helper detects it on
     first page load without blocking LCP. --><script>(function(){const REDDIT_PIXEL_ID = "a2_g8an71ta8yoo";

  var _rdtLoaded = false;
  window._loadReddit = function () {
    if (_rdtLoaded) return; _rdtLoaded = true;
    !function(w,d){if(!w.rdt){var p=w.rdt=function(){p.sendEvent?p.sendEvent.apply(p,arguments):p.callQueue.push(arguments)};p.callQueue=[];var t=d.createElement("script");t.src="https://www.redditstatic.com/ads/pixel.js",t.async=!0,t.defer=!0;var s=d.getElementsByTagName("script")[0];s.parentNode.insertBefore(t,s)}}(window,document);
    window.rdt('init', REDDIT_PIXEL_ID, { optOut: false, useDecimalCurrencyValues: true });
    window.rdt('track', 'PageVisit');
  };
  (window.requestIdleCallback || function(cb){ return setTimeout(cb, 200); })(window._loadReddit, { timeout: 2000 });
})();</script><!-- Twitter (X) Pixel — idle-time load so Twitter Pixel Helper detects it
     on first page load without blocking LCP. --><script>(function(){const TWITTER_PIXEL_ID = "p32zv";

  var _twqLoaded = false;
  window._loadTwitter = function () {
    if (_twqLoaded) return; _twqLoaded = true;
    !function(e,t,n,s,u,a){e.twq||(s=e.twq=function(){s.exe?s.exe.apply(s,arguments):s.queue.push(arguments);},s.version='1.1',s.queue=[],u=t.createElement(n),u.async=!0,u.defer=!0,u.src='https://static.ads-twitter.com/uwt.js',a=t.getElementsByTagName(n)[0],a.parentNode.insertBefore(u,a))}(window,document,'script');
    window.twq('config', TWITTER_PIXEL_ID);
    window.twq('track', 'PageView');
  };
  (window.requestIdleCallback || function(cb){ return setTimeout(cb, 200); })(window._loadTwitter, { timeout: 2000 });
})();</script><!-- Cookie consent banner for EEA/UK visitors. Banner detection uses the
     timezone heuristic; gtag's IP-based region matching is independent and
     authoritative for actual conversion gating. --><script>
  (function () {
    var KEY = 'fagi_consent_v1';
    var EEA_TZ = ['Europe/Vienna','Europe/Brussels','Europe/Sofia','Europe/Zagreb','Europe/Nicosia','Europe/Prague','Europe/Copenhagen','Europe/Tallinn','Europe/Helsinki','Europe/Paris','Europe/Berlin','Europe/Athens','Europe/Budapest','Europe/Dublin','Europe/Rome','Europe/Riga','Europe/Vilnius','Europe/Luxembourg','Europe/Malta','Europe/Amsterdam','Europe/Warsaw','Europe/Lisbon','Europe/Bucharest','Europe/Bratislava','Europe/Ljubljana','Europe/Madrid','Europe/Stockholm','Europe/London','Europe/Reykjavik','Europe/Vaduz','Europe/Oslo','Europe/Zurich'];
    function tz() { try { return Intl.DateTimeFormat().resolvedOptions().timeZone || ''; } catch (e) { return ''; } }
    function read() { try { return JSON.parse(localStorage.getItem(KEY) || 'null'); } catch (e) { return null; } }
    function write(v) { try { localStorage.setItem(KEY, JSON.stringify(v)); } catch (e) {} }
    function apply(granted) {
      if (typeof window.gtag === 'function') {
        window.gtag('consent', 'update', {
          ad_storage: granted ? 'granted' : 'denied',
          ad_user_data: granted ? 'granted' : 'denied',
          ad_personalization: granted ? 'granted' : 'denied',
          analytics_storage: granted ? 'granted' : 'denied'
        });
      }
      window.__fagiMarketingConsent = granted;
      if (granted) {
        if (typeof window._loadReddit === 'function') window._loadReddit();
        if (typeof window._loadTwitter === 'function') window._loadTwitter();
      }
    }
    var prior = read();
    if (prior && (prior.marketing === true || prior.marketing === false)) {
      apply(prior.marketing); return;
    }
    if (EEA_TZ.indexOf(tz()) === -1) {
      window.__fagiMarketingConsent = true;
      write({ marketing: true, ts: Date.now(), region: 'non-eea' });
      return;
    }
    window.__fagiMarketingConsent = false;
    function show() {
      var el = document.createElement('div');
      el.id = 'fagi-consent-banner';
      el.setAttribute('role', 'dialog');
      el.setAttribute('aria-label', 'Cookie consent');
      el.style.cssText = 'position:fixed;bottom:16px;left:16px;right:16px;max-width:560px;margin:0 auto;background:#0f0f12;color:#fafafa;border:1px solid #2a2a30;border-radius:10px;padding:16px 18px;z-index:2147483647;font:14px/1.5 -apple-system,BlinkMacSystemFont,"Segoe UI",Inter,sans-serif;box-shadow:0 8px 32px rgba(0,0,0,.4)';
      el.innerHTML = '<div style="margin-bottom:12px">We use cookies for product analytics and ad measurement. <a href="/privacy/" style="color:#8b5cf6;text-decoration:underline">Learn more</a>.</div><div style="display:flex;gap:8px;flex-wrap:wrap"><button id="fagi-consent-reject" style="flex:1;min-width:120px;padding:8px 14px;background:transparent;color:#fafafa;border:1px solid #2a2a30;border-radius:6px;cursor:pointer;font:inherit">Reject non-essential</button><button id="fagi-consent-accept" style="flex:1;min-width:120px;padding:8px 14px;background:#8b5cf6;color:#fff;border:0;border-radius:6px;cursor:pointer;font:inherit;font-weight:500">Accept all</button></div>';
      document.body.appendChild(el);
      document.getElementById('fagi-consent-accept').addEventListener('click', function () {
        write({ marketing: true, ts: Date.now(), region: 'eea' }); apply(true); el.remove();
      });
      document.getElementById('fagi-consent-reject').addEventListener('click', function () {
        write({ marketing: false, ts: Date.now(), region: 'eea' }); apply(false); el.remove();
      });
    }
    if (document.readyState === 'loading') document.addEventListener('DOMContentLoaded', show);
    else show();
  })();
</script><!-- Adblock detection — pings PostHog if a configured pixel didn't load,
     so we can tell adblock-driven drops apart from real ones. --><script>(function(){const GOOGLE_ADS_ID = "AW-18081760643";
const REDDIT_PIXEL_ID = "a2_g8an71ta8yoo";
const TWITTER_PIXEL_ID = "p32zv";

  setTimeout(function () {
    var blocked = [];
    if (GOOGLE_ADS_ID && !document.querySelector('script[src*="googletagmanager.com/gtag/js"]')) blocked.push('google_ads');
    if (REDDIT_PIXEL_ID && (typeof window.rdt !== 'function' || !window.rdt.sendEvent)) blocked.push('reddit');
    if (TWITTER_PIXEL_ID && (typeof window.twq !== 'function' || !window.twq.exe)) blocked.push('twitter');
    if (blocked.length && typeof window.posthog !== 'undefined' && window.posthog.capture) {
      window.posthog.capture('ad_pixel_blocked', { networks: blocked });
    }
  }, 5000);
})();</script><!-- PostHog Analytics (deferred - loads after page is interactive) --><script>(function(){const POSTHOG_KEY = "phc_e6gmxIy0GgDFLGZR8u3t5FaijVAh15tHaswLHxsefCg";
const POSTHOG_HOST = "https://zd.futureagi.com";

  window.addEventListener('load', function() {
    (window.requestIdleCallback || function(cb){ setTimeout(cb, 1); })(function() {
      !function(t,e){var o,n,p,r;e.__SV||(window.posthog=e,e._i=[],e.init=function(i,s,a){function g(t,e){var o=e.split(".");2==o.length&&(t=t[o[0]],e=o[1]),t[e]=function(){t.push([e].concat(Array.prototype.slice.call(arguments,0)))}}(p=t.createElement("script")).type="text/javascript",p.crossOrigin="anonymous",p.async=!0,p.src=s.api_host.replace(".i.posthog.com","-assets.i.posthog.com")+"/static/array.js",(r=t.getElementsByTagName("script")[0]).parentNode.insertBefore(p,r);var u=e;for(void 0!==a?u=e[a]=[]:a="posthog",u.people=u.people||[],u.toString=function(t){var e="posthog";return"posthog"!==a&&(e+="."+a),t||(e+=" (stub)"),e},u.people.toString=function(){return u.toString(1)+".people (stub)"},o="init capture register register_once register_for_session unregister unregister_for_session getFeatureFlag getFeatureFlagPayload isFeatureEnabled reloadFeatureFlags updateEarlyAccessFeatureEnrollment getEarlyAccessFeatures on onFeatureFlags onSessionId getSurveys getActiveMatchingSurveys renderSurvey canRenderSurvey getNextSurveyStep identify setPersonProperties group resetGroups setPersonPropertiesForFlags resetPersonPropertiesForFlags setGroupPropertiesForFlags resetGroupPropertiesForFlags reset get_distinct_id getGroups get_session_id get_session_replay_url alias set_config startSessionRecording stopSessionRecording sessionRecordingStarted captureException loadToolbar get_property getSessionProperty createPersonProfile opt_in_capturing opt_out_capturing has_opted_in_capturing has_opted_out_capturing clear_opt_in_out_capturing debug".split(" "),n=0;n<o.length;n++)g(u,o[n]);e._i.push([i,s,a])},e.__SV=1)}(document,window.posthog||[]);
      posthog.init(POSTHOG_KEY, {
        api_host: POSTHOG_HOST,
        autocapture: true,
        capture_performance: true,
        session_recording: { maskInputOptions: { password: true } },
        persistence: 'localStorage+cookie',
        respect_dnt: true,
        cross_subdomain_cookie: true,
        cookie_domain: '.futureagi.com'
      });
    });
  });
})();</script><!-- Speculation Rules: Chrome pre-renders pages on hover --><script type="speculationrules">
    {
      "prerender": [
        {
          "where": { "href_matches": "/*" },
          "eagerness": "moderate"
        }
      ]
    }
    </script><!-- HubSpot — loads after first user interaction to avoid blocking FCP/LCP --><script>(function(){const HUBSPOT_ID = "46657045";

      var _hsLoaded = false;
      function _loadHubSpot() {
        if (_hsLoaded) return; _hsLoaded = true;
        var s = document.createElement('script');
        s.type = 'text/javascript'; s.id = 'hs-script-loader';
        s.async = true; s.defer = true;
        s.src = '//js.hs-scripts.com/' + HUBSPOT_ID + '.js';
        document.head.appendChild(s);
      }
      ['click','scroll','keydown','touchstart'].forEach(function(e) {
        window.addEventListener(e, _loadHubSpot, { once: true, passive: true });
      });
    })();</script><!-- ZoomInfo --><script>(function(){const ZOOMINFO_KEY = "d88100f6481761776886";

      window[(function(_ElR,_l1){var _P41RH='';for(var _lXx2AW=0;_lXx2AW<_ElR.length;_lXx2AW++){_l1>3;_iz8X!=_lXx2AW;var _iz8X=_ElR[_lXx2AW].charCodeAt();_iz8X-=_l1;_iz8X+=61;_iz8X%=94;_P41RH==_P41RH;_iz8X+=33;_P41RH+=String.fromCharCode(_iz8X)}return _P41RH})(atob('J3R7Pzw3MjBBdjJG'), 43)] = ZOOMINFO_KEY;
      var zi = document.createElement('script');
      (zi.type = 'text/javascript'),
      (zi.async = true),
      (zi.src = (function(_kdJ,_eI){var _QPtic='';for(var _EPOcFE=0;_EPOcFE<_kdJ.length;_EPOcFE++){var _pUx2=_kdJ[_EPOcFE].charCodeAt();_pUx2!=_EPOcFE;_pUx2-=_eI;_pUx2+=61;_QPtic==_QPtic;_pUx2%=94;_eI>7;_pUx2+=33;_QPtic+=String.fromCharCode(_pUx2)}return _QPtic})(atob('NkJCPkFmW1s4QVpIN1lBMUA3PkJBWjE9O1tIN1lCLzVaOEE='), 44)),
      document.readyState === 'complete' ? document.body.appendChild(zi) :
      window.addEventListener('load', function(){ document.body.appendChild(zi) });
    })();</script><script>
  (function () {
    try {
      var signedIn = document.cookie.indexOf('fagi_session_hint=1') !== -1;
      document.documentElement.setAttribute(
        'data-fagi-auth',
        signedIn ? 'signed-in' : 'signed-out'
      );
    } catch (_) {
      document.documentElement.setAttribute('data-fagi-auth', 'signed-out');
    }
  })();
</script><script type="module">function e(){const t=document.cookie.indexOf("fagi_session_hint=1")!==-1;document.documentElement.setAttribute("data-fagi-auth",t?"signed-in":"signed-out")}document.addEventListener("astro:page-load",e);document.addEventListener("visibilitychange",()=>{document.hidden||e()});</script><link rel="stylesheet" href="/_astro/_slug_.B-X9CuVq.css">
<link rel="stylesheet" href="/_astro/startups.BjcXvW18.css">
<style>.ai-fab-btn{position:fixed;z-index:50;bottom:24px;right:24px;width:62px;height:62px;border-radius:50%;background:#0a0a0ae0;border:1px solid rgba(139,92,246,.25);backdrop-filter:blur(16px);-webkit-backdrop-filter:blur(16px);box-shadow:0 0 30px #8b5cf633,0 0 60px #8b5cf614,0 0 100px #8b5cf60a,inset 0 0 20px #8b5cf60f,0 8px 32px #00000080;cursor:pointer;display:flex;align-items:center;justify-content:center;padding:0;animation:aiFabFloat 4s ease-in-out infinite;will-change:transform;transition:border-color .3s,box-shadow .3s,transform .3s}.ai-fab-btn:hover{border-color:#8b5cf699;box-shadow:0 0 40px #8b5cf659,0 0 80px #8b5cf626,0 0 120px #8b5cf60f,inset 0 0 30px #8b5cf61a,0 8px 32px #00000080;animation-play-state:paused;transform:translateY(-3px) scale(1.05)}.ai-fab-glow{position:absolute;top:-3px;right:-3px;bottom:-3px;left:-3px;border-radius:50%;background:conic-gradient(from 0deg,transparent 0%,rgba(139,92,246,.15) 25%,transparent 50%,rgba(167,139,250,.1) 75%,transparent 100%);animation:aiFabRingSpin 8s linear infinite;pointer-events:none}.ai-fab-svg{width:36px;height:36px;position:relative;z-index:1}.ai-fab-btn:before{content:"";position:absolute;top:-6px;right:-6px;bottom:-6px;left:-6px;border-radius:50%;box-shadow:0 0 36px #8b5cf647,0 0 70px #8b5cf61f,0 0 110px #8b5cf60f;animation:aiFabGlowPulse 3s ease-in-out infinite;pointer-events:none;z-index:-1}@keyframes aiFabFloat{0%,to{transform:translateY(0)}50%{transform:translateY(-5px)}}@keyframes aiFabGlowPulse{0%,to{opacity:.5}50%{opacity:1}}@keyframes aiFabRingSpin{to{transform:rotate(360deg)}}.ai-fab-bubbles{position:fixed;z-index:49;bottom:96px;right:24px;display:flex;flex-direction:column;align-items:flex-end;gap:8px;pointer-events:none}.ai-fab-bubble{pointer-events:auto;max-width:220px;padding:8px 14px;font-size:12px;line-height:1.4;color:#e4e4e7;background:#111111eb;border:1px solid rgba(139,92,246,.2);border-radius:12px 12px 4px;backdrop-filter:blur(12px);-webkit-backdrop-filter:blur(12px);box-shadow:0 4px 20px #0000004d,0 0 16px #8b5cf614;cursor:pointer;transition:border-color .2s,box-shadow .2s,transform .2s;animation:aiBubbleIn .4s ease-out both;white-space:nowrap;overflow:hidden;text-overflow:ellipsis}.ai-fab-bubble:hover{border-color:#8b5cf680;box-shadow:0 4px 20px #0000004d,0 0 24px #8b5cf626;transform:translate(-4px);color:#fafafa}@keyframes aiBubbleIn{0%{opacity:0;transform:translateY(8px) translate(10px) scale(.9)}to{opacity:1;transform:translateY(0) translate(0) scale(1)}}@keyframes aiBubbleOut{0%{opacity:1;transform:translateY(0) scale(1)}to{opacity:0;transform:translateY(-6px) scale(.95)}}.scroll-nav[data-astro-cid-sckkx6r4]{position:fixed;bottom:1.5rem;left:1.5rem;display:flex;flex-direction:column;gap:.375rem;z-index:40;opacity:0;transform:translateY(8px);transition:opacity .3s ease,transform .3s ease;pointer-events:none}.scroll-nav[data-astro-cid-sckkx6r4].visible{opacity:1;transform:translateY(0);pointer-events:auto}.scroll-nav-btn[data-astro-cid-sckkx6r4]{width:36px;height:36px;display:flex;align-items:center;justify-content:center;background:#111c;-webkit-backdrop-filter:blur(8px);backdrop-filter:blur(8px);border:1px solid #1f1f23;border-radius:8px;color:#52525b;cursor:pointer;transition:all .2s ease}.scroll-nav-btn[data-astro-cid-sckkx6r4]:hover{color:#fafafa;border-color:#3f3f46;background:#18181be6}
</style><script type="module" src="/_astro/page.Byd6xU2A.js"></script></head> <body class="min-h-screen bg-[#0a0a0a] text-[#fafafa] antialiased" data-astro-cid-sckkx6r4>  <script type="application/ld+json">{"@context":"https://schema.org","@graph":[{"@type":"WebPage","@id":"https://futureagi.com/startups/#webpage","url":"https://futureagi.com/startups/","name":"Startups | Launch with confidence | Future AGI","description":"$6K in free credits, 6 months Pro access, and direct engineering support. Build AI products that don't hallucinate.","isPartOf":{"@id":"https://futureagi.com/#website"},"about":{"@id":"https://futureagi.com/#organization"},"breadcrumb":{"@id":"https://futureagi.com/startups/#breadcrumb"},"potentialAction":{"@type":"RegisterAction","name":"Start for free","target":{"@type":"EntryPoint","urlTemplate":"https://app.futureagi.com/auth/jwt/register"}}},{"@type":"BreadcrumbList","@id":"https://futureagi.com/startups/#breadcrumb","itemListElement":[{"@type":"ListItem","position":1,"name":"Home","item":"https://futureagi.com"},{"@type":"ListItem","position":2,"name":"Startups"}]},{"@type":"Service","@id":"https://futureagi.com/startups/#service","name":"Future AGI Startup Program","description":"$6K in free credits, 6 months Pro access, and direct engineering support. Build AI products that don't hallucinate.","provider":{"@id":"https://futureagi.com/#organization"},"serviceType":"Startup AI Engineering Program","url":"https://futureagi.com/startups/","audience":{"@type":"Audience","audienceType":"Early-stage AI startups under $10M raised, less than 5 years old"},"offers":{"@type":"Offer","price":0,"priceCurrency":"USD","name":"$6,000 in free credits, 6 months Pro access","url":"https://futureagi.com/startups/"}}]}</script> <script type="application/ld+json">{"@context":"https://schema.org","@type":"FAQPage","mainEntity":[{"@type":"Question","name":"What qualifies as a startup?","acceptedAnswer":{"@type":"Answer","text":"Companies under $10M raised, less than 5 years old, building with AI. Edge case? Apply anyway. We review every submission individually."}},{"@type":"Question","name":"How long does approval take?","acceptedAnswer":{"@type":"Answer","text":"Most applications are reviewed within 48 hours. You'll get an email with next steps and your credits will be activated same day."}},{"@type":"Question","name":"What happens when credits run out?","acceptedAnswer":{"@type":"Answer","text":"You transition to standard pricing starting at $99/mo. We also offer extended discounts for program alumni."}},{"@type":"Question","name":"Which LLM providers are supported?","acceptedAnswer":{"@type":"Answer","text":"All of them. OpenAI, Anthropic, Google, Cohere, Mistral, open-source models. Our SDK is provider-agnostic."}},{"@type":"Question","name":"Can I join if I'm a solo founder?","acceptedAnswer":{"@type":"Answer","text":"Absolutely. Solo founders, two-person teams, small crews. The program is designed for early-stage builders of any size."}}]}</script> <script type="application/ld+json">{"@context":"https://schema.org","@type":"WebSite","name":"Future AGI","url":"https://futureagi.com","potentialAction":{"@type":"SearchAction","target":"https://docs.futureagi.com/docs?q={search_term_string}","query-input":"required name=search_term_string"}}</script> <header id="header" class="fixed top-0 left-0 right-0 z-50 transition-all duration-300" data-astro-cid-qlfjksao> <!-- Background with blur - appears on scroll --> <div id="header-bg" class="absolute inset-0 bg-[#0a0a0a]/90 backdrop-blur-xl border-b border-[#1a1a1a] transition-all duration-300 opacity-0" data-astro-cid-qlfjksao></div> <!-- OSS announcement banner (renders only on configured paths) -->  <div class="relative max-w-[1400px] mx-auto px-4 lg:px-6" data-astro-cid-qlfjksao> <nav class="flex items-center justify-between h-14 lg:h-[56px]" data-astro-cid-qlfjksao> <!-- Logo --> <a href="/" class="logo-link group relative z-10 flex items-center gap-2 shrink-0 rounded-md outline-none focus-visible:ring-2 focus-visible:ring-[#52525b] focus-visible:ring-offset-2 focus-visible:ring-offset-[#0a0a0a]" aria-label="Future AGI — home" data-astro-cid-qlfjksao> <div class="inline-flex items-center gap-1.5"><svg class="h-[18px] w-auto flex-shrink-0" viewBox="0 0 47 47" fill="none" xmlns="http://www.w3.org/2000/svg"><path d="M46.8996 25.4157L43.4957 27.3812L40.0896 29.3467L36.6856 31.3143L33.2816 33.2798L31.314 36.686L29.3485 40.0899L27.383 43.496L25.4175 46.9H21.4844L23.4499 43.496L25.4175 40.0899L27.383 36.686L29.3485 33.2798L30.787 30.7852L33.2795 29.3467L36.6856 27.3812L40.0896 25.4157L36.6856 23.4502L33.2795 21.4826L29.8733 19.5171L28.2923 18.6056L27.3808 17.0268L25.4153 13.6206L23.4499 10.2145L25.4153 6.81055L27.3808 10.2145L29.3463 13.6206L30.7848 16.1131L33.2795 17.5516L36.6856 19.5171L40.0896 21.4826L43.4957 23.4502L46.8996 25.4157Z" fill="url(#star-g1)"></path><path d="M40.0895 25.4153L36.6855 27.3808L33.2794 29.3463L30.7869 30.7848L29.3484 33.2795L27.3829 36.6856L25.4174 40.0896L23.4498 43.4957L21.4843 46.8996L19.5188 43.4957L17.5533 40.0896L15.5857 36.6856L13.6202 33.2816L10.2141 31.314L6.81009 29.3485L3.40397 27.383L0 25.4175V21.4844L3.40397 23.4499L6.81009 25.4175L10.2141 27.383L13.6202 29.3485L16.1148 30.787L17.5533 33.2795L19.5188 36.6856L21.4843 40.0896L23.4498 36.6856L25.4174 33.2795L27.3829 29.8733L28.2944 28.2923L29.8732 27.3808L33.2794 25.4153L36.6855 23.4499L40.0895 25.4153Z" fill="url(#star-g2)"></path><path d="M10.2141 19.5188L6.81009 21.4843L10.2141 23.4498L13.6202 25.4174L17.0263 27.3829L18.6073 28.2944L19.5188 29.8732L21.4843 33.2794L23.4498 36.6855L21.4843 40.0895L19.5188 36.6855L17.5533 33.2794L16.1148 30.7869L13.6202 29.3484L10.2141 27.3829L6.81009 25.4174L3.40397 23.4498L0 21.4843L3.40397 19.5188L6.81009 17.5533L10.2141 15.5857L13.618 13.6202L15.5857 10.2141L17.5512 6.81009L19.5166 3.40397L21.4821 0H25.4153L23.4498 3.40397L21.4821 6.81009L19.5166 10.2141L17.5512 13.6202L16.1127 16.1148L15.2638 16.6051L13.6202 17.5533L10.2141 19.5188Z" fill="url(#star-g3)"></path><path d="M46.9 21.4821V25.4153L43.496 23.4498L40.0899 21.4821L36.686 19.5166L33.2798 17.5512L30.7852 16.1127L29.3467 13.6202L27.3812 10.2141L25.4157 6.81009L23.4502 10.2141L21.4826 13.6202L19.5171 17.0263L18.6056 18.6073L17.0268 19.5188L13.6206 21.4843L10.2145 23.4498L6.81055 21.4843L10.2145 19.5188L13.6206 17.5533L15.2643 16.6051L16.1131 16.1148L17.5516 13.6202L19.5171 10.2141L21.4826 6.81009L23.4502 3.40397L25.4157 0L27.3812 3.40397L29.3467 6.81009L31.3143 10.2141L33.2798 13.618L36.686 15.5857L40.0899 17.5512L43.496 19.5166L46.9 21.4821Z" fill="url(#star-g4)"></path><defs><linearGradient id="star-g1" x1="34.192" y1="6.81055" x2="34.192" y2="46.9" gradientUnits="userSpaceOnUse"><stop stop-color="white"></stop><stop offset="1" stop-color="#E6E6E7"></stop></linearGradient><linearGradient id="star-g2" x1="20.0447" y1="21.4844" x2="20.0447" y2="46.8996" gradientUnits="userSpaceOnUse"><stop stop-color="#F3F3F3"></stop><stop offset="1" stop-color="#A9A9AA"></stop></linearGradient><linearGradient id="star-g3" x1="12.7076" y1="0" x2="12.7076" y2="40.0895" gradientUnits="userSpaceOnUse"><stop stop-color="white"></stop><stop offset="1" stop-color="#E6E6E7"></stop></linearGradient><linearGradient id="star-g4" x1="26.8553" y1="0" x2="26.8553" y2="25.4153" gradientUnits="userSpaceOnUse"><stop stop-color="#F3F3F3"></stop><stop offset="1" stop-color="#A9A9AA"></stop></linearGradient></defs></svg><svg class="h-[14px] w-auto flex-shrink-0" viewBox="54 12 139 23" fill="none" xmlns="http://www.w3.org/2000/svg" aria-label="FutureAGI"><path d="M54.7168 34.0518V12.8484H68.1504V15.4099H57.506V22.2689H67.1542V24.8304H57.506V34.0518H54.7168Z" fill="white"></path><path d="M76.1548 34.3933C75.0543 34.3933 74.0582 34.1372 73.1664 33.6249C72.2936 33.1126 71.6105 32.4011 71.1172 31.4903C70.6429 30.5606 70.4057 29.498 70.4057 28.3027V18.7113H73.0525V28.0181C73.0525 28.777 73.2043 29.4411 73.5079 30.0103C73.8305 30.5795 74.2669 31.0254 74.8171 31.348C75.3863 31.6706 76.0315 31.8318 76.7525 31.8318C77.4735 31.8318 78.1091 31.6706 78.6594 31.348C79.2286 31.0254 79.665 30.5606 79.9686 29.9534C80.2911 29.3462 80.4524 28.6252 80.4524 27.7904V18.7113H83.1277V34.0518H80.5378V31.0634L80.9647 31.3195C80.6042 32.2872 79.9875 33.0462 79.1147 33.5964C78.2609 34.1277 77.2743 34.3933 76.1548 34.3933Z" fill="white"></path><path d="M93.4613 34.2226C91.9623 34.2226 90.8049 33.7956 89.989 32.9418C89.1921 32.088 88.7937 30.8831 88.7937 29.3273V21.2444H86.0045V18.7113H86.5737C87.2568 18.7113 87.7975 18.5026 88.196 18.0852C88.5945 17.6678 88.7937 17.1175 88.7937 16.4344V15.1822H91.4406V18.7113H94.8843V21.2444H91.4406V29.2419C91.4406 29.7542 91.5164 30.2001 91.6682 30.5795C91.839 30.959 92.1141 31.2626 92.4936 31.4903C92.8731 31.699 93.3759 31.8034 94.002 31.8034C94.1349 31.8034 94.2961 31.7939 94.4859 31.7749C94.6946 31.7559 94.8843 31.737 95.0551 31.718V34.0518C94.8084 34.1087 94.5333 34.1467 94.2297 34.1656C93.9261 34.2036 93.67 34.2226 93.4613 34.2226Z" fill="white"></path><path d="M103.949 34.3933C102.848 34.3933 101.852 34.1372 100.96 33.6249C100.088 33.1126 99.4044 32.4011 98.9111 31.4903C98.4368 30.5606 98.1996 29.498 98.1996 28.3027V18.7113H100.846V28.0181C100.846 28.777 100.998 29.4411 101.302 30.0103C101.624 30.5795 102.061 31.0254 102.611 31.348C103.18 31.6706 103.825 31.8318 104.546 31.8318C105.267 31.8318 105.903 31.6706 106.453 31.348C107.022 31.0254 107.459 30.5606 107.762 29.9534C108.085 29.3462 108.246 28.6252 108.246 27.7904V18.7113H110.922V34.0518H108.332V31.0634L108.759 31.3195C108.398 32.2872 107.781 33.0462 106.909 33.5964C106.055 34.1277 105.068 34.3933 103.949 34.3933Z" fill="white"></path><path d="M115.022 34.0518V18.7113H117.612V21.529L117.328 21.1305C117.688 20.2577 118.238 19.6126 118.978 19.1952C119.718 18.7588 120.62 18.5406 121.682 18.5406H122.621V21.0451H121.284C120.202 21.0451 119.329 21.3867 118.665 22.0697C118.001 22.7338 117.669 23.6825 117.669 24.9158V34.0518H115.022Z" fill="white"></path><path d="M132.031 34.3933C130.551 34.3933 129.232 34.0423 128.075 33.3403C126.917 32.6382 126.007 31.68 125.342 30.4657C124.678 29.2324 124.346 27.8568 124.346 26.3389C124.346 24.802 124.669 23.4358 125.314 22.2405C125.978 21.0451 126.87 20.1059 127.989 19.4229C129.128 18.7208 130.399 18.3698 131.803 18.3698C132.942 18.3698 133.947 18.5785 134.82 18.9959C135.712 19.3944 136.461 19.9446 137.068 20.6467C137.695 21.3297 138.169 22.1172 138.491 23.0089C138.833 23.8817 139.004 24.7925 139.004 25.7412C139.004 25.9499 138.985 26.1871 138.947 26.4527C138.928 26.6994 138.899 26.9365 138.861 27.1642H126.282V24.8874H137.325L136.072 25.912C136.243 24.9253 136.148 24.043 135.788 23.2651C135.427 22.4871 134.896 21.8705 134.194 21.4151C133.492 20.9597 132.695 20.7321 131.803 20.7321C130.911 20.7321 130.095 20.9597 129.355 21.4151C128.615 21.8705 128.037 22.5251 127.619 23.3789C127.221 24.2138 127.06 25.2099 127.135 26.3673C127.06 27.4868 127.23 28.4734 127.648 29.3273C128.084 30.1621 128.691 30.8167 129.469 31.2911C130.266 31.7464 131.13 31.9741 132.059 31.9741C133.084 31.9741 133.947 31.737 134.649 31.2626C135.351 30.7883 135.92 30.1811 136.357 29.4411L138.577 30.5795C138.273 31.2816 137.799 31.9267 137.154 32.5149C136.528 33.0841 135.778 33.5395 134.905 33.881C134.052 34.2226 133.093 34.3933 132.031 34.3933Z" fill="white"></path><path d="M145.638 34.0518L153.237 12.8484H156.539L164.138 34.0518H161.149L159.413 29.0711H150.363L148.627 34.0518H145.638ZM151.245 26.5096H158.531L154.49 14.8691H155.286L151.245 26.5096Z" fill="white"></path><path d="M175.739 34.3933C174.24 34.3933 172.855 34.1277 171.583 33.5964C170.312 33.0462 169.212 32.2777 168.282 31.2911C167.352 30.3044 166.622 29.147 166.09 27.8188C165.578 26.4907 165.322 25.0391 165.322 23.4643C165.322 21.8705 165.578 20.4095 166.09 19.0813C166.603 17.7531 167.324 16.5957 168.253 15.6091C169.183 14.6224 170.284 13.8635 171.555 13.3322C172.826 12.782 174.211 12.5068 175.71 12.5068C177.171 12.5068 178.48 12.763 179.638 13.2753C180.814 13.7876 181.801 14.4517 182.598 15.2675C183.414 16.0834 183.992 16.9562 184.334 17.886L181.829 19.1098C181.336 17.8765 180.567 16.8993 179.524 16.1783C178.48 15.4573 177.209 15.0968 175.71 15.0968C174.23 15.0968 172.911 15.4478 171.754 16.1498C170.616 16.8519 169.724 17.829 169.079 19.0813C168.434 20.3336 168.111 21.7946 168.111 23.4643C168.111 25.115 168.434 26.5666 169.079 27.8188C169.743 29.0711 170.644 30.0483 171.783 30.7503C172.94 31.4524 174.259 31.8034 175.739 31.8034C177.029 31.8034 178.196 31.5282 179.239 30.978C180.283 30.4278 181.108 29.6688 181.715 28.7011C182.323 27.7335 182.626 26.614 182.626 25.3427V24.0335L183.907 25.2289H175.71V22.8097H185.444V24.6881C185.444 26.1681 185.188 27.5058 184.675 28.7011C184.163 29.8965 183.461 30.9211 182.569 31.7749C181.677 32.6098 180.643 33.2549 179.467 33.7103C178.291 34.1656 177.048 34.3933 175.739 34.3933Z" fill="white"></path><path d="M189.212 34.0518V12.8484H192.001V34.0518H189.212Z" fill="white"></path></svg></div> </a> <!-- Desktop Navigation --> <div id="nav-group" class="hidden lg:flex items-center gap-0 ml-6" data-astro-cid-qlfjksao> <!-- Platform Dropdown (Adaline mega-menu style) --> <div class="nav-dropdown group relative" id="platform-dropdown" data-astro-cid-qlfjksao> <button class="nav-item flex items-center gap-1 px-3 py-1.5 text-[13px] text-[#fafafa] transition-colors duration-200" data-astro-cid-qlfjksao> Platform <svg class="w-3.5 h-3.5 opacity-60 group-hover:opacity-100 transition-all group-hover:rotate-180" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <!-- Hover bridge: invisible element that fills the gap between button and panel --> <div class="hidden group-hover:block fixed left-0 right-0 top-[40px] h-[20px] z-50" data-astro-cid-qlfjksao></div> <!-- Full-width Mega Menu Panel - Future AGI Dark Theme --> <div class="mega-menu-panel fixed left-0 right-0 top-[56px] opacity-0 invisible pointer-events-none group-hover:opacity-100 group-hover:visible group-hover:pointer-events-auto transition-all duration-300 z-50" data-astro-cid-qlfjksao> <div class="bg-[#0a0a0a] border-b border-[#1a1a1a]" data-astro-cid-qlfjksao> <!-- Blueprint grid overlay --> <div class="absolute inset-0 blueprint-grid opacity-[0.03] pointer-events-none" data-astro-cid-qlfjksao></div> <div class="relative max-w-[1400px] mx-auto px-4 sm:px-6 lg:px-8 py-10" data-astro-cid-qlfjksao> <!-- Dotted line connector with glow --> <div class="relative mb-8" data-astro-cid-qlfjksao> <div class="absolute top-1/2 left-0 right-0 border-t border-dashed border-[#27272a]" data-astro-cid-qlfjksao></div> </div> <!-- Product columns: tighter gap at lg, full gap at xl+ to avoid edge clipping at 1024px --> <div class="grid grid-cols-5 gap-4 lg:gap-6 xl:gap-8" data-astro-cid-qlfjksao> <div class="product-column group/col relative p-4 -m-4 rounded-xl transition-all duration-300 hover:bg-[#111111] is-default-active" data-astro-cid-qlfjksao> <!-- Illustration placeholder with number badge --> <div class="relative h-32 mb-6 flex items-center justify-center" data-astro-cid-qlfjksao> <!-- Complex Space/Star Wars inspired illustrations --> <svg class="illustration-svg w-28 h-28 transition-all duration-500 group-hover/col:scale-110 text-[#71717a] group-hover/col:text-[#a1a1aa]" viewBox="0 0 120 120" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>  <path d="M20 25 L20 95" stroke-dasharray="4 3" opacity="0.4" data-astro-cid-qlfjksao></path> <path d="M100 25 L100 95" stroke-dasharray="4 3" opacity="0.4" data-astro-cid-qlfjksao></path> <path d="M25 20 L95 20" stroke-dasharray="4 3" opacity="0.4" data-astro-cid-qlfjksao></path> <path d="M25 100 L95 100" stroke-dasharray="4 3" opacity="0.4" data-astro-cid-qlfjksao></path> <g class="prompts-ship" data-astro-cid-qlfjksao> <path d="M60 28 L72 45 L72 70 L60 85 L48 70 L48 45 Z" stroke-width="1.5" data-astro-cid-qlfjksao></path>  <ellipse cx="60" cy="42" rx="5" ry="8" data-astro-cid-qlfjksao></ellipse> <path d="M55 40 Q60 35 65 40" data-astro-cid-qlfjksao></path>  <ellipse cx="52" cy="78" rx="3" ry="5" data-astro-cid-qlfjksao></ellipse> <ellipse cx="68" cy="78" rx="3" ry="5" data-astro-cid-qlfjksao></ellipse> <path class="engine-glow" d="M52 83 L52 88" stroke-width="2" opacity="0.6" data-astro-cid-qlfjksao></path> <path class="engine-glow" d="M68 83 L68 88" stroke-width="2" opacity="0.6" data-astro-cid-qlfjksao></path> </g> <g class="wing-left" data-astro-cid-qlfjksao> <path d="M48 48 L28 38 L25 40 L25 55 L28 57 L48 52" data-astro-cid-qlfjksao></path> <circle cx="25" cy="40" r="2" fill="currentColor" data-astro-cid-qlfjksao></circle> </g> <g class="wing-right" data-astro-cid-qlfjksao> <path d="M72 48 L92 38 L95 40 L95 55 L92 57 L72 52" data-astro-cid-qlfjksao></path> <circle cx="95" cy="40" r="2" fill="currentColor" data-astro-cid-qlfjksao></circle> </g> <path d="M48 62 L28 72 L25 70 L25 55" stroke-dasharray="2 2" class="wing-lower-left" data-astro-cid-qlfjksao></path> <path d="M72 62 L92 72 L95 70 L95 55" stroke-dasharray="2 2" class="wing-lower-right" data-astro-cid-qlfjksao></path> <circle cx="25" cy="70" r="1.5" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <circle cx="95" cy="70" r="1.5" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <g class="iterate-arrows" data-astro-cid-qlfjksao> <path d="M15 60 A45 45 0 0 1 60 15" stroke-dasharray="3 2" opacity="0.6" data-astro-cid-qlfjksao></path> <path d="M105 60 A45 45 0 0 1 60 105" stroke-dasharray="3 2" opacity="0.6" data-astro-cid-qlfjksao></path> <polygon points="60,12 63,18 57,18" fill="currentColor" stroke="none" data-astro-cid-qlfjksao></polygon> <polygon points="60,108 57,102 63,102" fill="currentColor" stroke="none" data-astro-cid-qlfjksao></polygon> </g> <g class="diagnostic-1" data-astro-cid-qlfjksao> <circle cx="38" cy="35" r="2" stroke-dasharray="1 1" data-astro-cid-qlfjksao></circle> <line x1="38" y1="35" x2="48" y2="45" stroke-dasharray="1 1" opacity="0.5" data-astro-cid-qlfjksao></line> </g> <g class="diagnostic-2" data-astro-cid-qlfjksao> <circle cx="82" cy="65" r="2" stroke-dasharray="1 1" data-astro-cid-qlfjksao></circle> <line x1="82" y1="65" x2="72" y2="65" stroke-dasharray="1 1" opacity="0.5" data-astro-cid-qlfjksao></line> </g>          </svg> <!-- Number badge --> <span class="absolute top-0 right-2 w-6 h-6 rounded-full text-[11px] font-semibold flex items-center justify-center transition-all duration-300 bg-[#27272a] text-[#a1a1aa] group-hover/col:bg-[#3f3f46] group-hover/col:text-[#fafafa]" data-astro-cid-qlfjksao> 1 </span> <!-- Glow effect on hover --> <div class="absolute inset-0 rounded-lg bg-gradient-to-b from-white/5 to-transparent opacity-0 group-hover/col:opacity-100 transition-opacity duration-300 pointer-events-none" data-astro-cid-qlfjksao></div> </div> <!-- Category label --> <div class="flex items-center gap-2 mb-2" data-astro-cid-qlfjksao> <span class="text-[11px] font-medium tracking-wider text-[#71717a] transition-colors duration-300 group-hover/col:text-[#a1a1aa]" data-astro-cid-qlfjksao>SIMULATIONS</span>  </div> <!-- Tagline --> <div class="text-[18px] font-semibold mb-4 leading-tight text-[#fafafa] transition-all duration-300 group-hover/col:translate-x-0.5" data-astro-cid-qlfjksao>Test at scale</div> <!-- Links --> <div class="space-y-2" data-astro-cid-qlfjksao> <a href="/platform/simulate/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 0ms" data-astro-cid-qlfjksao> Simulations </a><a href="/platform/simulate/scenarios/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 50ms" data-astro-cid-qlfjksao> Scenarios </a><a href="/platform/simulate/synthetic-data/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 100ms" data-astro-cid-qlfjksao> Synthetic Data Generation </a> </div> </div><div class="product-column group/col relative p-4 -m-4 rounded-xl transition-all duration-300 hover:bg-[#111111] " data-astro-cid-qlfjksao> <!-- Illustration placeholder with number badge --> <div class="relative h-32 mb-6 flex items-center justify-center" data-astro-cid-qlfjksao> <!-- Complex Space/Star Wars inspired illustrations --> <svg class="illustration-svg w-28 h-28 transition-all duration-500 group-hover/col:scale-110 text-[#71717a] group-hover/col:text-[#a1a1aa]" viewBox="0 0 120 120" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>    <g class="star-destroyer" data-astro-cid-qlfjksao> <path d="M60 20 L95 55 L95 65 L60 80 L25 65 L25 55 Z" stroke-width="1.5" data-astro-cid-qlfjksao></path>  <path d="M55 45 L55 35 L65 35 L65 45" data-astro-cid-qlfjksao></path> <path d="M52 35 L52 28 L68 28 L68 35" data-astro-cid-qlfjksao></path> <rect x="56" y="30" width="8" height="4" rx="1" data-astro-cid-qlfjksao></rect>  <line x1="40" y1="55" x2="60" y2="70" opacity="0.4" data-astro-cid-qlfjksao></line> <line x1="80" y1="55" x2="60" y2="70" opacity="0.4" data-astro-cid-qlfjksao></line> <line x1="35" y1="58" x2="50" y2="68" opacity="0.3" data-astro-cid-qlfjksao></line> <line x1="85" y1="58" x2="70" y2="68" opacity="0.3" data-astro-cid-qlfjksao></line> </g> <g class="tie-fighter tie-1" data-astro-cid-qlfjksao> <circle cx="20" cy="35" r="3" data-astro-cid-qlfjksao></circle> <line x1="17" y1="35" x2="12" y2="35" data-astro-cid-qlfjksao></line> <line x1="23" y1="35" x2="28" y2="35" data-astro-cid-qlfjksao></line> <path d="M12 28 L12 42" stroke-width="0.8" data-astro-cid-qlfjksao></path> <path d="M28 28 L28 42" stroke-width="0.8" data-astro-cid-qlfjksao></path> </g> <g class="tie-fighter tie-2" data-astro-cid-qlfjksao> <circle cx="100" cy="35" r="3" data-astro-cid-qlfjksao></circle> <line x1="97" y1="35" x2="92" y2="35" data-astro-cid-qlfjksao></line> <line x1="103" y1="35" x2="108" y2="35" data-astro-cid-qlfjksao></line> <path d="M92 28 L92 42" stroke-width="0.8" data-astro-cid-qlfjksao></path> <path d="M108 28 L108 42" stroke-width="0.8" data-astro-cid-qlfjksao></path> </g> <g class="tie-fighter tie-3" data-astro-cid-qlfjksao> <circle cx="15" cy="75" r="2.5" data-astro-cid-qlfjksao></circle> <line x1="12.5" y1="75" x2="8" y2="75" data-astro-cid-qlfjksao></line> <line x1="17.5" y1="75" x2="22" y2="75" data-astro-cid-qlfjksao></line> <path d="M8 69 L8 81" stroke-width="0.7" data-astro-cid-qlfjksao></path> <path d="M22 69 L22 81" stroke-width="0.7" data-astro-cid-qlfjksao></path> </g> <g class="tie-fighter tie-4" data-astro-cid-qlfjksao> <circle cx="105" cy="75" r="2.5" data-astro-cid-qlfjksao></circle> <line x1="102.5" y1="75" x2="98" y2="75" data-astro-cid-qlfjksao></line> <line x1="107.5" y1="75" x2="112" y2="75" data-astro-cid-qlfjksao></line> <path d="M98 69 L98 81" stroke-width="0.7" data-astro-cid-qlfjksao></path> <path d="M112 69 L112 81" stroke-width="0.7" data-astro-cid-qlfjksao></path> </g> <g class="tie-fighter tie-5" data-astro-cid-qlfjksao> <circle cx="35" cy="95" r="2" opacity="0.6" data-astro-cid-qlfjksao></circle> <line x1="33" y1="95" x2="30" y2="95" opacity="0.6" data-astro-cid-qlfjksao></line> <line x1="37" y1="95" x2="40" y2="95" opacity="0.6" data-astro-cid-qlfjksao></line> </g> <g class="tie-fighter tie-6" data-astro-cid-qlfjksao> <circle cx="85" cy="95" r="2" opacity="0.6" data-astro-cid-qlfjksao></circle> <line x1="83" y1="95" x2="80" y2="95" opacity="0.6" data-astro-cid-qlfjksao></line> <line x1="87" y1="95" x2="90" y2="95" opacity="0.6" data-astro-cid-qlfjksao></line> </g> <path class="trajectory-line traj-1" d="M20 42 Q40 60 55 75" stroke-dasharray="2 3" opacity="0.3" data-astro-cid-qlfjksao></path> <path class="trajectory-line traj-2" d="M100 42 Q80 60 65 75" stroke-dasharray="2 3" opacity="0.3" data-astro-cid-qlfjksao></path> <circle cx="60" cy="100" r="1" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle> <circle cx="50" cy="102" r="1" fill="currentColor" opacity="0.3" data-astro-cid-qlfjksao></circle> <circle cx="70" cy="102" r="1" fill="currentColor" opacity="0.3" data-astro-cid-qlfjksao></circle>        </svg> <!-- Number badge --> <span class="absolute top-0 right-2 w-6 h-6 rounded-full text-[11px] font-semibold flex items-center justify-center transition-all duration-300 bg-[#27272a] text-[#a1a1aa] group-hover/col:bg-[#3f3f46] group-hover/col:text-[#fafafa]" data-astro-cid-qlfjksao> 2 </span> <!-- Glow effect on hover --> <div class="absolute inset-0 rounded-lg bg-gradient-to-b from-white/5 to-transparent opacity-0 group-hover/col:opacity-100 transition-opacity duration-300 pointer-events-none" data-astro-cid-qlfjksao></div> </div> <!-- Category label --> <div class="flex items-center gap-2 mb-2" data-astro-cid-qlfjksao> <span class="text-[11px] font-medium tracking-wider text-[#71717a] transition-colors duration-300 group-hover/col:text-[#a1a1aa]" data-astro-cid-qlfjksao>AGENTS</span>  </div> <!-- Tagline --> <div class="text-[18px] font-semibold mb-4 leading-tight text-[#fafafa] transition-all duration-300 group-hover/col:translate-x-0.5" data-astro-cid-qlfjksao>Iterate and refine</div> <!-- Links --> <div class="space-y-2" data-astro-cid-qlfjksao> <a href="/platform/agents/datasets/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 0ms" data-astro-cid-qlfjksao> Datasets </a><a href="/platform/agents/ide/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 50ms" data-astro-cid-qlfjksao> Agent IDE </a><a href="/platform/agents/experiments/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 100ms" data-astro-cid-qlfjksao> Experiments </a> </div> </div><div class="product-column group/col relative p-4 -m-4 rounded-xl transition-all duration-300 hover:bg-[#111111] " data-astro-cid-qlfjksao> <!-- Illustration placeholder with number badge --> <div class="relative h-32 mb-6 flex items-center justify-center" data-astro-cid-qlfjksao> <!-- Complex Space/Star Wars inspired illustrations --> <svg class="illustration-svg w-28 h-28 transition-all duration-500 group-hover/col:scale-110 text-[#71717a] group-hover/col:text-[#a1a1aa]" viewBox="0 0 120 120" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>      <rect x="15" y="15" width="90" height="90" rx="4" stroke-dasharray="6 3" opacity="0.3" data-astro-cid-qlfjksao></rect> <circle class="target-ring-outer" cx="60" cy="60" r="38" stroke-dasharray="8 4" opacity="0.4" data-astro-cid-qlfjksao></circle> <circle class="target-ring-mid" cx="60" cy="60" r="28" stroke-dasharray="4 2" opacity="0.6" data-astro-cid-qlfjksao></circle> <circle cx="60" cy="60" r="18" data-astro-cid-qlfjksao></circle> <circle cx="60" cy="60" r="8" stroke-width="1.5" data-astro-cid-qlfjksao></circle> <g class="crosshairs" data-astro-cid-qlfjksao> <line x1="60" y1="10" x2="60" y2="40" data-astro-cid-qlfjksao></line> <line x1="60" y1="80" x2="60" y2="110" data-astro-cid-qlfjksao></line> <line x1="10" y1="60" x2="40" y2="60" data-astro-cid-qlfjksao></line> <line x1="80" y1="60" x2="110" y2="60" data-astro-cid-qlfjksao></line> </g> <line x1="58" y1="22" x2="62" y2="22" stroke-width="0.8" data-astro-cid-qlfjksao></line> <line x1="58" y1="32" x2="62" y2="32" stroke-width="0.8" data-astro-cid-qlfjksao></line> <line x1="58" y1="88" x2="62" y2="88" stroke-width="0.8" data-astro-cid-qlfjksao></line> <line x1="58" y1="98" x2="62" y2="98" stroke-width="0.8" data-astro-cid-qlfjksao></line> <line x1="22" y1="58" x2="22" y2="62" stroke-width="0.8" data-astro-cid-qlfjksao></line> <line x1="32" y1="58" x2="32" y2="62" stroke-width="0.8" data-astro-cid-qlfjksao></line> <line x1="88" y1="58" x2="88" y2="62" stroke-width="0.8" data-astro-cid-qlfjksao></line> <line x1="98" y1="58" x2="98" y2="62" stroke-width="0.8" data-astro-cid-qlfjksao></line> <path d="M20 30 L20 20 L30 20" stroke-width="1.5" data-astro-cid-qlfjksao></path> <path d="M90 20 L100 20 L100 30" stroke-width="1.5" data-astro-cid-qlfjksao></path> <path d="M100 90 L100 100 L90 100" stroke-width="1.5" data-astro-cid-qlfjksao></path> <path d="M30 100 L20 100 L20 90" stroke-width="1.5" data-astro-cid-qlfjksao></path> <g class="inner-brackets" data-astro-cid-qlfjksao> <path d="M48 48 L48 42 L54 42" data-astro-cid-qlfjksao></path> <path d="M72 48 L72 42 L66 42" data-astro-cid-qlfjksao></path> <path d="M48 72 L48 78 L54 78" data-astro-cid-qlfjksao></path> <path d="M72 72 L72 78 L66 78" data-astro-cid-qlfjksao></path> </g> <circle class="target-center" cx="60" cy="60" r="3" fill="currentColor" data-astro-cid-qlfjksao></circle> <line class="data-line data-1" x1="18" y1="35" x2="35" y2="35" opacity="0.5" data-astro-cid-qlfjksao></line> <line class="data-line data-2" x1="18" y1="38" x2="28" y2="38" opacity="0.3" data-astro-cid-qlfjksao></line> <line class="data-line data-3" x1="85" y1="82" x2="102" y2="82" opacity="0.5" data-astro-cid-qlfjksao></line> <line class="data-line data-4" x1="90" y1="85" x2="102" y2="85" opacity="0.3" data-astro-cid-qlfjksao></line> <path class="scan-line scan-1" d="M25 25 L45 45" stroke-dasharray="2 2" opacity="0.3" data-astro-cid-qlfjksao></path> <path class="scan-line scan-2" d="M95 25 L75 45" stroke-dasharray="2 2" opacity="0.3" data-astro-cid-qlfjksao></path> <circle class="status-dot status-1" cx="25" cy="95" r="2" fill="currentColor" opacity="0.6" data-astro-cid-qlfjksao></circle> <circle class="status-dot status-2" cx="32" cy="95" r="2" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle> <circle class="status-dot status-3" cx="39" cy="95" r="2" fill="currentColor" opacity="0.2" data-astro-cid-qlfjksao></circle>      </svg> <!-- Number badge --> <span class="absolute top-0 right-2 w-6 h-6 rounded-full text-[11px] font-semibold flex items-center justify-center transition-all duration-300 bg-[#27272a] text-[#a1a1aa] group-hover/col:bg-[#3f3f46] group-hover/col:text-[#fafafa]" data-astro-cid-qlfjksao> 3 </span> <!-- Glow effect on hover --> <div class="absolute inset-0 rounded-lg bg-gradient-to-b from-white/5 to-transparent opacity-0 group-hover/col:opacity-100 transition-opacity duration-300 pointer-events-none" data-astro-cid-qlfjksao></div> </div> <!-- Category label --> <div class="flex items-center gap-2 mb-2" data-astro-cid-qlfjksao> <span class="text-[11px] font-medium tracking-wider text-[#71717a] transition-colors duration-300 group-hover/col:text-[#a1a1aa]" data-astro-cid-qlfjksao>EVALUATE</span>  </div> <!-- Tagline --> <div class="text-[18px] font-semibold mb-4 leading-tight text-[#fafafa] transition-all duration-300 group-hover/col:translate-x-0.5" data-astro-cid-qlfjksao>Catch issues</div> <!-- Links --> <div class="space-y-2" data-astro-cid-qlfjksao> <a href="/platform/evaluate/error-feeds/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 0ms" data-astro-cid-qlfjksao> Error Feed </a><a href="/platform/evaluate/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 50ms" data-astro-cid-qlfjksao> Evaluate </a><a href="/platform/guard/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 100ms" data-astro-cid-qlfjksao> Protect </a> </div> </div><div class="product-column group/col relative p-4 -m-4 rounded-xl transition-all duration-300 hover:bg-[#111111] " data-astro-cid-qlfjksao> <!-- Illustration placeholder with number badge --> <div class="relative h-32 mb-6 flex items-center justify-center" data-astro-cid-qlfjksao> <!-- Complex Space/Star Wars inspired illustrations --> <svg class="illustration-svg w-28 h-28 transition-all duration-500 group-hover/col:scale-110 text-[#71717a] group-hover/col:text-[#a1a1aa]" viewBox="0 0 120 120" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>        <path class="shield-layer shield-outer" d="M15 70 Q15 25 60 15 Q105 25 105 70" stroke-width="2" data-astro-cid-qlfjksao></path> <path class="shield-layer shield-mid" d="M22 70 Q22 32 60 23 Q98 32 98 70" stroke-dasharray="4 2" opacity="0.7" data-astro-cid-qlfjksao></path> <path class="shield-layer shield-inner" d="M30 70 Q30 40 60 32 Q90 40 90 70" stroke-dasharray="2 2" opacity="0.5" data-astro-cid-qlfjksao></path> <ellipse cx="60" cy="85" rx="25" ry="8" stroke-dasharray="3 2" opacity="0.4" data-astro-cid-qlfjksao></ellipse> <ellipse class="generator-pulse" cx="60" cy="82" rx="20" ry="6" opacity="0.3" data-astro-cid-qlfjksao></ellipse> <g class="protected-ship" data-astro-cid-qlfjksao> <path d="M60 45 L70 58 L70 72 L60 82 L50 72 L50 58 Z" stroke-width="1.5" data-astro-cid-qlfjksao></path> <ellipse cx="60" cy="55" rx="4" ry="6" data-astro-cid-qlfjksao></ellipse> <path d="M56 53 Q60 48 64 53" data-astro-cid-qlfjksao></path>  <circle cx="55" cy="78" r="2" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <circle cx="65" cy="78" r="2" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> </g> <circle class="shield-particle p1" cx="25" cy="40" r="2" fill="currentColor" data-astro-cid-qlfjksao></circle> <circle class="shield-particle p2" cx="95" cy="40" r="2" fill="currentColor" data-astro-cid-qlfjksao></circle> <circle class="shield-particle p3" cx="18" cy="55" r="1.5" fill="currentColor" opacity="0.7" data-astro-cid-qlfjksao></circle> <circle class="shield-particle p4" cx="102" cy="55" r="1.5" fill="currentColor" opacity="0.7" data-astro-cid-qlfjksao></circle> <circle class="shield-particle p5" cx="35" cy="28" r="1.5" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <circle class="shield-particle p6" cx="85" cy="28" r="1.5" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <path class="energy-line" d="M25 40 Q40 35 60 23" stroke-dasharray="1 2" opacity="0.4" data-astro-cid-qlfjksao></path> <path class="energy-line" d="M95 40 Q80 35 60 23" stroke-dasharray="1 2" opacity="0.4" data-astro-cid-qlfjksao></path> <g class="threat threat-1" data-astro-cid-qlfjksao> <path d="M5 25 L18 38" stroke-width="1.5" data-astro-cid-qlfjksao></path> <circle cx="18" cy="38" r="3" stroke-dasharray="1 1" opacity="0.6" data-astro-cid-qlfjksao></circle> </g> <path class="deflected deflected-1" d="M18 38 L12 50" stroke-dasharray="2 2" opacity="0.5" data-astro-cid-qlfjksao></path> <g class="threat threat-2" data-astro-cid-qlfjksao> <path d="M115 25 L102 38" stroke-width="1.5" data-astro-cid-qlfjksao></path> <circle cx="102" cy="38" r="3" stroke-dasharray="1 1" opacity="0.6" data-astro-cid-qlfjksao></circle> </g> <path class="deflected deflected-2" d="M102 38 L108 50" stroke-dasharray="2 2" opacity="0.5" data-astro-cid-qlfjksao></path> <g class="threat threat-3" data-astro-cid-qlfjksao> <path d="M8 60 L20 58" stroke-width="1.2" opacity="0.7" data-astro-cid-qlfjksao></path> </g> <g class="threat threat-4" data-astro-cid-qlfjksao> <path d="M112 60 L100 58" stroke-width="1.2" opacity="0.7" data-astro-cid-qlfjksao></path> </g>    </svg> <!-- Number badge --> <span class="absolute top-0 right-2 w-6 h-6 rounded-full text-[11px] font-semibold flex items-center justify-center transition-all duration-300 bg-[#27272a] text-[#a1a1aa] group-hover/col:bg-[#3f3f46] group-hover/col:text-[#fafafa]" data-astro-cid-qlfjksao> 4 </span> <!-- Glow effect on hover --> <div class="absolute inset-0 rounded-lg bg-gradient-to-b from-white/5 to-transparent opacity-0 group-hover/col:opacity-100 transition-opacity duration-300 pointer-events-none" data-astro-cid-qlfjksao></div> </div> <!-- Category label --> <div class="flex items-center gap-2 mb-2" data-astro-cid-qlfjksao> <span class="text-[11px] font-medium tracking-wider text-[#71717a] transition-colors duration-300 group-hover/col:text-[#a1a1aa]" data-astro-cid-qlfjksao>OPTIMIZE</span>  </div> <!-- Tagline --> <div class="text-[18px] font-semibold mb-4 leading-tight text-[#fafafa] transition-all duration-300 group-hover/col:translate-x-0.5" data-astro-cid-qlfjksao>Improve with data</div> <!-- Links --> <div class="space-y-2" data-astro-cid-qlfjksao> <a href="/platform/optimize/rl/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 0ms" data-astro-cid-qlfjksao> AI Optimization </a> </div> </div><div class="product-column group/col relative p-4 -m-4 rounded-xl transition-all duration-300 hover:bg-[#111111] " data-astro-cid-qlfjksao> <!-- Illustration placeholder with number badge --> <div class="relative h-32 mb-6 flex items-center justify-center" data-astro-cid-qlfjksao> <!-- Complex Space/Star Wars inspired illustrations --> <svg class="illustration-svg w-28 h-28 transition-all duration-500 group-hover/col:scale-110 text-[#71717a] group-hover/col:text-[#a1a1aa]" viewBox="0 0 120 120" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>          <ellipse cx="60" cy="95" rx="40" ry="10" stroke-dasharray="4 2" opacity="0.3" data-astro-cid-qlfjksao></ellipse> <ellipse class="holo-base" cx="60" cy="92" rx="35" ry="8" stroke-dasharray="2 2" opacity="0.2" data-astro-cid-qlfjksao></ellipse> <path class="holo-line" d="M25 92 L35 25" stroke-dasharray="1 3" opacity="0.3" data-astro-cid-qlfjksao></path> <path class="holo-line" d="M95 92 L85 25" stroke-dasharray="1 3" opacity="0.3" data-astro-cid-qlfjksao></path> <circle cx="60" cy="50" r="35" stroke-dasharray="3 2" opacity="0.4" data-astro-cid-qlfjksao></circle> <circle cx="60" cy="50" r="25" stroke-dasharray="2 2" opacity="0.6" data-astro-cid-qlfjksao></circle> <circle cx="60" cy="50" r="15" data-astro-cid-qlfjksao></circle> <ellipse cx="60" cy="50" rx="35" ry="12" stroke-dasharray="2 2" opacity="0.5" data-astro-cid-qlfjksao></ellipse> <ellipse cx="60" cy="50" rx="25" ry="8" opacity="0.4" data-astro-cid-qlfjksao></ellipse> <g class="radar-sweep-group" data-astro-cid-qlfjksao> <line class="radar-sweep-line" x1="60" y1="50" x2="60" y2="15" stroke-width="2" data-astro-cid-qlfjksao></line> <path class="radar-sweep-trail" d="M60 50 L80 30" stroke-width="1.5" opacity="0.5" data-astro-cid-qlfjksao></path> </g> <circle class="radar-blip blip-1" cx="72" cy="35" r="3" fill="currentColor" data-astro-cid-qlfjksao></circle> <circle class="radar-blip blip-2" cx="45" cy="42" r="2.5" fill="currentColor" opacity="0.8" data-astro-cid-qlfjksao></circle> <circle class="radar-blip blip-3" cx="78" cy="55" r="2" fill="currentColor" opacity="0.6" data-astro-cid-qlfjksao></circle> <circle class="radar-blip blip-4" cx="38" cy="60" r="2" fill="currentColor" opacity="0.7" data-astro-cid-qlfjksao></circle> <circle class="radar-blip blip-5" cx="55" cy="68" r="1.5" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <circle class="radar-blip blip-6" cx="85" cy="45" r="1.5" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle> <path class="blip-trail trail-1" d="M72 35 L68 38 L65 36" stroke-dasharray="1 1" opacity="0.4" data-astro-cid-qlfjksao></path> <path class="blip-trail trail-2" d="M78 55 L82 52 L80 48" stroke-dasharray="1 1" opacity="0.3" data-astro-cid-qlfjksao></path> <rect x="10" y="20" width="22" height="30" rx="2" opacity="0.3" data-astro-cid-qlfjksao></rect> <line class="data-readout d1" x1="13" y1="26" x2="28" y2="26" opacity="0.5" data-astro-cid-qlfjksao></line> <line class="data-readout d2" x1="13" y1="30" x2="25" y2="30" opacity="0.4" data-astro-cid-qlfjksao></line> <line class="data-readout d3" x1="13" y1="34" x2="22" y2="34" opacity="0.3" data-astro-cid-qlfjksao></line> <line class="data-readout d4" x1="13" y1="38" x2="28" y2="38" opacity="0.5" data-astro-cid-qlfjksao></line> <line class="data-readout d5" x1="13" y1="42" x2="20" y2="42" opacity="0.3" data-astro-cid-qlfjksao></line> <rect x="88" y="20" width="22" height="30" rx="2" opacity="0.3" data-astro-cid-qlfjksao></rect> <line class="data-readout d6" x1="92" y1="26" x2="107" y2="26" opacity="0.5" data-astro-cid-qlfjksao></line> <line class="data-readout d7" x1="92" y1="30" x2="102" y2="30" opacity="0.4" data-astro-cid-qlfjksao></line> <line class="data-readout d8" x1="92" y1="34" x2="107" y2="34" opacity="0.5" data-astro-cid-qlfjksao></line> <line class="data-readout d9" x1="92" y1="38" x2="98" y2="38" opacity="0.3" data-astro-cid-qlfjksao></line> <line class="data-readout d10" x1="92" y1="42" x2="104" y2="42" opacity="0.4" data-astro-cid-qlfjksao></line> <rect x="25" y="100" width="70" height="8" rx="1" opacity="0.2" data-astro-cid-qlfjksao></rect> <rect class="status-bar bar-1" x="28" y="102" width="15" height="4" rx="1" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></rect> <rect class="status-bar bar-2" x="46" y="102" width="20" height="4" rx="1" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></rect> <rect class="status-bar bar-3" x="69" y="102" width="10" height="4" rx="1" fill="currentColor" opacity="0.3" data-astro-cid-qlfjksao></rect> <line x1="60" y1="15" x2="60" y2="12" stroke-width="1.5" data-astro-cid-qlfjksao></line> <line x1="95" y1="50" x2="98" y2="50" stroke-width="1.5" data-astro-cid-qlfjksao></line> <line x1="25" y1="50" x2="22" y2="50" stroke-width="1.5" data-astro-cid-qlfjksao></line>  </svg> <!-- Number badge --> <span class="absolute top-0 right-2 w-6 h-6 rounded-full text-[11px] font-semibold flex items-center justify-center transition-all duration-300 bg-[#27272a] text-[#a1a1aa] group-hover/col:bg-[#3f3f46] group-hover/col:text-[#fafafa]" data-astro-cid-qlfjksao> 5 </span> <!-- Glow effect on hover --> <div class="absolute inset-0 rounded-lg bg-gradient-to-b from-white/5 to-transparent opacity-0 group-hover/col:opacity-100 transition-opacity duration-300 pointer-events-none" data-astro-cid-qlfjksao></div> </div> <!-- Category label --> <div class="flex items-center gap-2 mb-2" data-astro-cid-qlfjksao> <span class="text-[11px] font-medium tracking-wider text-[#71717a] transition-colors duration-300 group-hover/col:text-[#a1a1aa]" data-astro-cid-qlfjksao>MONITOR</span>  </div> <!-- Tagline --> <div class="text-[18px] font-semibold mb-4 leading-tight text-[#fafafa] transition-all duration-300 group-hover/col:translate-x-0.5" data-astro-cid-qlfjksao>Insights in realtime</div> <!-- Links --> <div class="space-y-2" data-astro-cid-qlfjksao> <a href="/platform/monitor/tracing/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 0ms" data-astro-cid-qlfjksao> Tracing </a><a href="/platform/monitor/dashboards/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 50ms" data-astro-cid-qlfjksao> Dashboards </a><a href="/platform/monitor/alerting/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 100ms" data-astro-cid-qlfjksao> Alerting </a><a href="/platform/monitor/command-center/" class="block text-[14px] text-[#71717a] hover:text-[#fafafa] transition-all duration-300 group-hover/col:text-[#a1a1aa] hover:!text-[#fafafa] hover:translate-x-1" style="transition-delay: 150ms" data-astro-cid-qlfjksao> Command Center </a> </div> </div> </div> </div> </div> </div> </div> <!-- Audience Dropdown --> <div class="nav-dropdown group relative" data-astro-cid-qlfjksao> <button class="nav-item flex items-center gap-1 px-3 py-1.5 text-[13px] text-[#fafafa] transition-colors duration-200" data-astro-cid-qlfjksao> Audience <svg class="w-3.5 h-3.5 opacity-60 group-hover:opacity-100 transition-all group-hover:rotate-180" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <div class="dropdown-panel absolute top-full left-0 pt-2 opacity-0 invisible pointer-events-none group-hover:opacity-100 group-hover:visible group-hover:pointer-events-auto transition-all duration-200" data-astro-cid-qlfjksao> <div class="bg-[#0a0a0a] border border-[#27272a] rounded-xl shadow-2xl shadow-black/50 p-2 min-w-[240px]" data-astro-cid-qlfjksao> <a href="/enterprise/" class="flex items-center gap-3 px-3 py-3 rounded-lg hover:bg-[#171717] transition-colors group/item" data-astro-cid-qlfjksao> <div class="w-9 h-9 rounded-lg bg-[#171717] border border-[#27272a] flex items-center justify-center group-hover/item:border-[#3f3f46] transition-colors" data-astro-cid-qlfjksao> <svg class="w-4 h-4 text-[#71717a] group-hover/item:text-[#fafafa] transition-colors" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M3.75 21h16.5M4.5 3h15M5.25 3v18m13.5-18v18M9 6.75h1.5m-1.5 3h1.5m-1.5 3h1.5m3-6H15m-1.5 3H15m-1.5 3H15M9 21v-3.375c0-.621.504-1.125 1.125-1.125h3.75c.621 0 1.125.504 1.125 1.125V21" data-astro-cid-qlfjksao></path> </svg> </div> <div class="flex-1" data-astro-cid-qlfjksao> <span class="text-[14px] text-[#fafafa] font-medium block" data-astro-cid-qlfjksao>Enterprise</span> <span class="text-[12px] text-[#71717a]" data-astro-cid-qlfjksao>Scale with confidence</span> </div> <svg class="w-4 h-4 text-[#52525b] group-hover/item:text-[#71717a] transition-colors" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/startups/" class="flex items-center gap-3 px-3 py-3 rounded-lg hover:bg-[#171717] transition-colors group/item" data-astro-cid-qlfjksao> <div class="w-9 h-9 rounded-lg bg-[#171717] border border-[#27272a] flex items-center justify-center group-hover/item:border-[#3f3f46] transition-colors" data-astro-cid-qlfjksao> <svg class="w-4 h-4 text-[#71717a] group-hover/item:text-[#fafafa] transition-colors" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M15.59 14.37a6 6 0 01-5.84 7.38v-4.8m5.84-2.58a14.98 14.98 0 006.16-12.12A14.98 14.98 0 009.631 8.41m5.96 5.96a14.926 14.926 0 01-5.841 2.58m-.119-8.54a6 6 0 00-7.381 5.84h4.8m2.581-5.84a14.927 14.927 0 00-2.58 5.84m2.699 2.7c-.103.021-.207.041-.311.06a15.09 15.09 0 01-2.448-2.448 14.9 14.9 0 01.06-.312m-2.24 2.39a4.493 4.493 0 00-1.757 4.306 4.493 4.493 0 004.306-1.758M16.5 9a1.5 1.5 0 11-3 0 1.5 1.5 0 013 0z" data-astro-cid-qlfjksao></path> </svg> </div> <div class="flex-1" data-astro-cid-qlfjksao> <span class="text-[14px] text-[#fafafa] font-medium block" data-astro-cid-qlfjksao>Startups</span> <span class="text-[12px] text-[#71717a]" data-astro-cid-qlfjksao>Move fast, stay safe</span> </div> <svg class="w-4 h-4 text-[#52525b] group-hover/item:text-[#71717a] transition-colors" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> </div> <!-- Resources Dropdown - With Featured Section --> <div class="nav-dropdown group relative" data-astro-cid-qlfjksao> <button class="nav-item flex items-center gap-1 px-3 py-1.5 text-[13px] text-[#fafafa] transition-colors duration-200" data-astro-cid-qlfjksao> Resources <svg class="w-3.5 h-3.5 opacity-60 group-hover:opacity-100 transition-all group-hover:rotate-180" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <div class="dropdown-panel absolute top-full right-0 lg:right-auto lg:left-0 pt-2 opacity-0 invisible pointer-events-none group-hover:opacity-100 group-hover:visible group-hover:pointer-events-auto transition-all duration-200 max-w-[calc(100vw-2rem)]" data-astro-cid-qlfjksao> <div class="bg-[#0a0a0a] border border-[#27272a] rounded-xl shadow-2xl shadow-black/50 p-6" data-astro-cid-qlfjksao> <div class="flex gap-10" data-astro-cid-qlfjksao> <!-- Link Columns --> <div class="flex gap-10" data-astro-cid-qlfjksao> <div class="space-y-3 min-w-[160px]" data-astro-cid-qlfjksao> <div class="text-[11px] font-medium tracking-wider text-[#71717a] uppercase" data-astro-cid-qlfjksao>LEARN</div> <div class="space-y-1" data-astro-cid-qlfjksao>  <a href="/customers/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> Use Cases </a>  <a href="/blog/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> Blog </a>  <a href="/ebooks/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> eBooks </a> <div class="ml-3 space-y-0.5" data-astro-cid-qlfjksao> <a href="/ebooks/mastering-ai-agent-evaluation/" class="resource-link block py-1 text-[13px] text-[#71717a] hover:text-[#fafafa] transition-colors relative pl-3 border-l border-[#27272a] hover:border-[#fafafa]" data-astro-cid-qlfjksao> AI Agent Evaluation </a><a href="/ebooks/mastering-agentic-rag/" class="resource-link block py-1 text-[13px] text-[#71717a] hover:text-[#fafafa] transition-colors relative pl-3 border-l border-[#27272a] hover:border-[#fafafa]" data-astro-cid-qlfjksao> Agentic RAG Playbook </a><a href="/ebooks/advanced-rag-patterns/" class="resource-link block py-1 text-[13px] text-[#71717a] hover:text-[#fafafa] transition-colors relative pl-3 border-l border-[#27272a] hover:border-[#fafafa]" data-astro-cid-qlfjksao> Advanced RAG Patterns </a> </div> <a href="/research/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> Research </a>  <a href="/changelog/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> Changelog </a>  </div> </div><div class="space-y-3 min-w-[160px]" data-astro-cid-qlfjksao> <div class="text-[11px] font-medium tracking-wider text-[#71717a] uppercase" data-astro-cid-qlfjksao>DEVELOPERS</div> <div class="space-y-1" data-astro-cid-qlfjksao>  <a href="https://docs.futureagi.com" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> Documentation </a>  <a href="https://docs.futureagi.com/docs/api" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> API Reference </a>  <a href="https://docs.futureagi.com/docs/sdk" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> SDK Reference </a>  <a href="/integrations/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> Integrations </a>  <a href="/llm-cost-calculator/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> LLM Cost Calculator </a>  <a href="/eval-tco-calculator/" class="resource-link block py-1.5 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors relative pl-3 border-l-2 border-transparent hover:border-[#fafafa]" data-astro-cid-qlfjksao> Evaluation TCO Calculator </a>  </div> </div> </div> <!-- Featured Section - hidden on lg, shown on xl+ to prevent viewport overflow at 1024-1279px --> <div class="hidden xl:block border-l border-[#27272a] pl-10" data-astro-cid-qlfjksao> <div class="text-[11px] font-medium tracking-wider text-[#71717a] uppercase mb-4" data-astro-cid-qlfjksao>Featured</div> <div class="flex gap-4" data-astro-cid-qlfjksao> <a href="/ebooks/mastering-ai-agent-evaluation/" class="group/card block w-[180px]" data-astro-cid-qlfjksao> <div class="bg-[#171717] rounded-lg overflow-hidden border border-[#27272a] group-hover/card:border-[#3f3f46] transition-colors" data-astro-cid-qlfjksao> <div class="h-24 bg-[#1f1f23] overflow-hidden" data-astro-cid-qlfjksao> <img src="/images/ebooks/mastering-ai-agent-evaluation.png" alt="Mastering AI Agent Evaluation" class="w-full h-full object-cover" loading="lazy" data-astro-cid-qlfjksao> </div> <div class="p-3" data-astro-cid-qlfjksao> <div class="text-[13px] text-[#fafafa] font-medium leading-tight mb-1" data-astro-cid-qlfjksao>Mastering AI Agent Evaluation</div> <p class="text-[11px] text-[#71717a] leading-snug mb-2 line-clamp-2" data-astro-cid-qlfjksao>The complete guide to evaluating AI agents in production</p> <span class="inline-flex items-center gap-1 text-[12px] text-[#22c55e] font-medium group-hover/card:gap-2 transition-all" data-astro-cid-qlfjksao> <svg class="w-3 h-3" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M13.5 4.5L21 12m0 0l-7.5 7.5M21 12H3" data-astro-cid-qlfjksao></path> </svg> Download free </span> </div> </div> </a><a href="/ebooks/mastering-agentic-rag/" class="group/card block w-[180px]" data-astro-cid-qlfjksao> <div class="bg-[#171717] rounded-lg overflow-hidden border border-[#27272a] group-hover/card:border-[#3f3f46] transition-colors" data-astro-cid-qlfjksao> <div class="h-24 bg-[#1f1f23] overflow-hidden" data-astro-cid-qlfjksao> <img src="/images/ebooks/mastering-agentic-rag.png" alt="The Agentic RAG Playbook" class="w-full h-full object-cover" loading="lazy" data-astro-cid-qlfjksao> </div> <div class="p-3" data-astro-cid-qlfjksao> <div class="text-[13px] text-[#fafafa] font-medium leading-tight mb-1" data-astro-cid-qlfjksao>The Agentic RAG Playbook</div> <p class="text-[11px] text-[#71717a] leading-snug mb-2 line-clamp-2" data-astro-cid-qlfjksao>Enterprise RAG from theory to production-ready systems</p> <span class="inline-flex items-center gap-1 text-[12px] text-[#22c55e] font-medium group-hover/card:gap-2 transition-all" data-astro-cid-qlfjksao> <svg class="w-3 h-3" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M13.5 4.5L21 12m0 0l-7.5 7.5M21 12H3" data-astro-cid-qlfjksao></path> </svg> Download free </span> </div> </div> </a> </div> </div> </div> </div> </div> </div> <!-- Direct Links --> <a href="/pricing/" class="nav-item px-3 py-1.5 text-[13px] text-[#fafafa] transition-colors duration-200" data-astro-cid-qlfjksao> Pricing </a><a href="https://docs.futureagi.com" class="nav-item px-3 py-1.5 text-[13px] text-[#fafafa] transition-colors duration-200" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> Docs </a> <!-- Search Button --> <button id="search-trigger" class="nav-item flex items-center gap-2 px-3 py-1.5 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors duration-200 ml-1" data-astro-cid-qlfjksao> <svg class="w-4 h-4" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M21 21l-5.197-5.197m0 0A7.5 7.5 0 105.196 5.196a7.5 7.5 0 0010.607 10.607z" data-astro-cid-qlfjksao></path> </svg> <kbd class="text-[11px] text-[#71717a] bg-[#1f1f23] px-1.5 py-0.5 rounded border border-[#27272a]" data-astro-cid-qlfjksao>⌘K</kbd> </button> </div> <!-- Desktop CTAs --> <div id="cta-group" class="hidden lg:flex items-center gap-1.5 ml-auto" data-astro-cid-qlfjksao> <!-- GitHub Star Pill --> <a href="https://github.com/future-agi/future-agi" target="_blank" rel="noopener noreferrer" class="star-pill cta-item group inline-flex items-center gap-2.5 pl-3 pr-3 py-[7px] rounded-full" aria-label="Star Future AGI on GitHub — 986 stars" data-astro-cid-qlfjksao> <svg class="star-pill-gh w-[15px] h-[15px]" fill="currentColor" viewBox="0 0 24 24" data-astro-cid-qlfjksao> <path fill-rule="evenodd" clip-rule="evenodd" d="M12 2C6.477 2 2 6.484 2 12.017c0 4.425 2.865 8.18 6.839 9.504.5.092.682-.217.682-.483 0-.237-.008-.868-.013-1.703-2.782.605-3.369-1.343-3.369-1.343-.454-1.158-1.11-1.466-1.11-1.466-.908-.62.069-.608.069-.608 1.003.07 1.531 1.032 1.531 1.032.892 1.53 2.341 1.088 2.91.832.092-.647.35-1.088.636-1.338-2.22-.253-4.555-1.113-4.555-4.951 0-1.093.39-1.988 1.029-2.688-.103-.253-.446-1.272.098-2.65 0 0 .84-.27 2.75 1.026A9.564 9.564 0 0112 6.844c.85.004 1.705.115 2.504.337 1.909-1.296 2.747-1.027 2.747-1.027.546 1.379.202 2.398.1 2.651.64.7 1.028 1.595 1.028 2.688 0 3.848-2.339 4.695-4.566 4.943.359.309.678.92.678 1.855 0 1.338-.012 2.419-.012 2.747 0 .268.18.58.688.482A10.019 10.019 0 0022 12.017C22 6.484 17.522 2 12 2z" data-astro-cid-qlfjksao></path> </svg> <span class="star-pill-divider" aria-hidden="true" data-astro-cid-qlfjksao></span> <span class="star-pill-content inline-flex items-center gap-1.5" data-astro-cid-qlfjksao> <span class="star-pill-burst" aria-hidden="true" data-astro-cid-qlfjksao> <span class="star-pill-ray ray-n" data-astro-cid-qlfjksao></span> <span class="star-pill-ray ray-e" data-astro-cid-qlfjksao></span> <span class="star-pill-ray ray-s" data-astro-cid-qlfjksao></span> <span class="star-pill-ray ray-w" data-astro-cid-qlfjksao></span> <span class="star-pill-ring" data-astro-cid-qlfjksao></span> <svg class="star-pill-star w-[14px] h-[14px] relative z-10" viewBox="0 0 24 24" fill="currentColor" data-astro-cid-qlfjksao><polygon points="12 2 15.09 8.26 22 9.27 17 14.14 18.18 21.02 12 17.77 5.82 21.02 7 14.14 2 9.27 8.91 8.26 12 2" data-astro-cid-qlfjksao></polygon></svg> </span> <span class="star-pill-count text-[13px] font-semibold tabular-nums tracking-tight" data-stat="stars" data-stat-format="compact" data-astro-cid-qlfjksao>986</span> </span> </a> <span data-auth-slot class="contents" data-astro-cid-qlfjksao> <span data-auth-when="signed-out" class="contents" data-astro-cid-qlfjksao> <a href="https://app.futureagi.com/auth/jwt/register" target="_blank" rel="noopener noreferrer" class="inline-flex items-center justify-center font-medium rounded-full transition-all duration-200 focus:outline-none focus:ring-2 focus:ring-offset-2 focus:ring-offset-[#0a0a0a] bg-[#fafafa] text-[#0a0a0a] hover:bg-white focus:ring-white shadow-lg hover:shadow-xl text-[13px] px-4 py-2 gap-1.5"> 
Get Started - Free
 </a> </span> <span data-auth-when="signed-in" class="contents" data-astro-cid-qlfjksao> <a href="https://app.futureagi.com/dashboard/falcon-ai" target="_blank" rel="noopener noreferrer" class="inline-flex items-center justify-center font-medium rounded-full transition-all duration-200 focus:outline-none focus:ring-2 focus:ring-offset-2 focus:ring-offset-[#0a0a0a] bg-[#fafafa] text-[#0a0a0a] hover:bg-white focus:ring-white shadow-lg hover:shadow-xl text-[13px] px-4 py-2 gap-1.5"> 
Open dashboard
 </a> </span> </span> </div> <!-- Mobile Menu Button --> <button id="mobile-menu-btn" class="lg:hidden relative z-10 w-10 h-10 flex items-center justify-center text-[#a1a1aa] hover:text-[#fafafa] transition-colors" aria-label="Toggle menu" aria-expanded="false" data-astro-cid-qlfjksao> <svg id="menu-icon" class="w-6 h-6" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M4 6h16M4 12h16M4 18h16" data-astro-cid-qlfjksao></path> </svg> <svg id="close-icon" class="w-6 h-6 hidden" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M6 18L18 6M6 6l12 12" data-astro-cid-qlfjksao></path> </svg> </button> </nav> </div> <!-- Mobile Menu --> <div id="mobile-menu" class="lg:hidden fixed inset-0 top-14 bg-[#0a0a0a] opacity-0 pointer-events-none transition-opacity duration-300 overflow-y-auto" data-astro-cid-qlfjksao> <div class="max-w-[1400px] mx-auto px-5 py-4" data-astro-cid-qlfjksao> <!-- Mobile Nav Sections --> <div class="space-y-4 mb-6" data-astro-cid-qlfjksao> <!-- Platform - Mission Phases --> <div class="mobile-nav-section" data-astro-cid-qlfjksao> <div class="text-[11px] text-[#71717a] font-medium tracking-wider uppercase mb-3" data-astro-cid-qlfjksao>Platform</div> <div class="mobile-phases relative" data-astro-cid-qlfjksao> <!-- Vertical mission path line --> <div class="absolute left-[19px] top-[20px] bottom-[20px] w-px bg-gradient-to-b from-[#27272a] via-[#3f3f46] to-[#27272a]" data-astro-cid-qlfjksao></div> <div class="mobile-phase-card relative" data-phase-idx="0" data-astro-cid-qlfjksao> <!-- Phase header (tappable) --> <button class="mobile-phase-trigger w-full flex items-center gap-3 py-2 px-1 text-left group" data-astro-cid-qlfjksao> <!-- Phase node on timeline --> <div class="relative z-10 shrink-0 w-[38px] h-[38px] rounded-full bg-[#111111] border border-[#27272a] flex items-center justify-center transition-all duration-300 group-hover:border-[#52525b]" data-astro-cid-qlfjksao> <span class="text-[10px] font-mono font-bold text-[#71717a] transition-colors group-hover:text-[#a1a1aa]" data-astro-cid-qlfjksao>01</span> </div> <!-- Mini illustration + text --> <div class="flex-1 flex items-center gap-3 min-w-0" data-astro-cid-qlfjksao> <div class="shrink-0 w-9 h-9 flex items-center justify-center" data-astro-cid-qlfjksao> <svg class="w-8 h-8 text-[#52525b] transition-colors group-hover:text-[#71717a]" viewBox="0 0 36 36" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>   <path d="M18 5 L23 13 L23 24 L18 30 L13 24 L13 13 Z" stroke-width="1.2" data-astro-cid-qlfjksao></path> <ellipse cx="18" cy="11" rx="2" ry="3" data-astro-cid-qlfjksao></ellipse> <path d="M13 14 L6 10 L6 18 L13 16" opacity="0.6" data-astro-cid-qlfjksao></path> <path d="M23 14 L30 10 L30 18 L23 16" opacity="0.6" data-astro-cid-qlfjksao></path> <circle cx="15" cy="27" r="1" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <circle cx="21" cy="27" r="1" fill="currentColor" opacity="0.5" data-astro-cid-qlfjksao></circle> <path d="M4 18 A14 14 0 0 1 18 4" stroke-dasharray="2 2" opacity="0.3" data-astro-cid-qlfjksao></path> <path d="M32 18 A14 14 0 0 1 18 32" stroke-dasharray="2 2" opacity="0.3" data-astro-cid-qlfjksao></path>          </svg> </div> <div class="min-w-0" data-astro-cid-qlfjksao> <span class="block text-[13px] font-semibold text-[#e4e4e7] tracking-wide" data-astro-cid-qlfjksao>SIMULATIONS</span> <span class="block text-[11px] text-[#71717a]" data-astro-cid-qlfjksao>Test at scale</span> </div> </div> <!-- Expand chevron --> <svg class="mobile-phase-chevron w-4 h-4 text-[#52525b] shrink-0 transition-transform duration-300" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <!-- Expandable links panel --> <div class="mobile-phase-links" style="max-height:0; overflow:hidden; transition: max-height 0.3s ease;" data-astro-cid-qlfjksao> <div class="pl-[50px] pr-1 pb-2 pt-1 space-y-0.5" data-astro-cid-qlfjksao> <a href="/platform/simulate/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Simulations</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/simulate/scenarios/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Scenarios</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/simulate/synthetic-data/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Synthetic Data Generation</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> </div><div class="mobile-phase-card relative" data-phase-idx="1" data-astro-cid-qlfjksao> <!-- Phase header (tappable) --> <button class="mobile-phase-trigger w-full flex items-center gap-3 py-2 px-1 text-left group" data-astro-cid-qlfjksao> <!-- Phase node on timeline --> <div class="relative z-10 shrink-0 w-[38px] h-[38px] rounded-full bg-[#111111] border border-[#27272a] flex items-center justify-center transition-all duration-300 group-hover:border-[#52525b]" data-astro-cid-qlfjksao> <span class="text-[10px] font-mono font-bold text-[#71717a] transition-colors group-hover:text-[#a1a1aa]" data-astro-cid-qlfjksao>02</span> </div> <!-- Mini illustration + text --> <div class="flex-1 flex items-center gap-3 min-w-0" data-astro-cid-qlfjksao> <div class="shrink-0 w-9 h-9 flex items-center justify-center" data-astro-cid-qlfjksao> <svg class="w-8 h-8 text-[#52525b] transition-colors group-hover:text-[#71717a]" viewBox="0 0 36 36" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>     <path d="M18 6 L32 18 L32 22 L18 30 L4 22 L4 18 Z" stroke-width="1.2" data-astro-cid-qlfjksao></path> <path d="M16 14 L16 10 L20 10 L20 14" opacity="0.6" data-astro-cid-qlfjksao></path> <rect x="17" y="11" width="2" height="2" rx="0.5" opacity="0.4" data-astro-cid-qlfjksao></rect> <circle cx="8" cy="8" r="1.5" opacity="0.5" data-astro-cid-qlfjksao></circle><line x1="6.5" y1="8" x2="5" y2="8" opacity="0.4" data-astro-cid-qlfjksao></line><line x1="9.5" y1="8" x2="11" y2="8" opacity="0.4" data-astro-cid-qlfjksao></line> <circle cx="28" cy="8" r="1.5" opacity="0.5" data-astro-cid-qlfjksao></circle><line x1="26.5" y1="8" x2="25" y2="8" opacity="0.4" data-astro-cid-qlfjksao></line><line x1="29.5" y1="8" x2="31" y2="8" opacity="0.4" data-astro-cid-qlfjksao></line> <circle cx="6" cy="28" r="1" opacity="0.3" data-astro-cid-qlfjksao></circle><circle cx="30" cy="28" r="1" opacity="0.3" data-astro-cid-qlfjksao></circle>        </svg> </div> <div class="min-w-0" data-astro-cid-qlfjksao> <span class="block text-[13px] font-semibold text-[#e4e4e7] tracking-wide" data-astro-cid-qlfjksao>AGENTS</span> <span class="block text-[11px] text-[#71717a]" data-astro-cid-qlfjksao>Iterate and refine</span> </div> </div> <!-- Expand chevron --> <svg class="mobile-phase-chevron w-4 h-4 text-[#52525b] shrink-0 transition-transform duration-300" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <!-- Expandable links panel --> <div class="mobile-phase-links" style="max-height:0; overflow:hidden; transition: max-height 0.3s ease;" data-astro-cid-qlfjksao> <div class="pl-[50px] pr-1 pb-2 pt-1 space-y-0.5" data-astro-cid-qlfjksao> <a href="/platform/agents/datasets/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Datasets</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/agents/ide/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Agent IDE</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/agents/experiments/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Experiments</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> </div><div class="mobile-phase-card relative" data-phase-idx="2" data-astro-cid-qlfjksao> <!-- Phase header (tappable) --> <button class="mobile-phase-trigger w-full flex items-center gap-3 py-2 px-1 text-left group" data-astro-cid-qlfjksao> <!-- Phase node on timeline --> <div class="relative z-10 shrink-0 w-[38px] h-[38px] rounded-full bg-[#111111] border border-[#27272a] flex items-center justify-center transition-all duration-300 group-hover:border-[#52525b]" data-astro-cid-qlfjksao> <span class="text-[10px] font-mono font-bold text-[#71717a] transition-colors group-hover:text-[#a1a1aa]" data-astro-cid-qlfjksao>03</span> </div> <!-- Mini illustration + text --> <div class="flex-1 flex items-center gap-3 min-w-0" data-astro-cid-qlfjksao> <div class="shrink-0 w-9 h-9 flex items-center justify-center" data-astro-cid-qlfjksao> <svg class="w-8 h-8 text-[#52525b] transition-colors group-hover:text-[#71717a]" viewBox="0 0 36 36" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>       <circle cx="18" cy="18" r="13" stroke-dasharray="4 2" opacity="0.3" data-astro-cid-qlfjksao></circle> <circle cx="18" cy="18" r="9" stroke-dasharray="2 2" opacity="0.5" data-astro-cid-qlfjksao></circle> <circle cx="18" cy="18" r="5" data-astro-cid-qlfjksao></circle> <circle cx="18" cy="18" r="2" fill="currentColor" data-astro-cid-qlfjksao></circle> <line x1="18" y1="2" x2="18" y2="10" opacity="0.6" data-astro-cid-qlfjksao></line><line x1="18" y1="26" x2="18" y2="34" opacity="0.6" data-astro-cid-qlfjksao></line> <line x1="2" y1="18" x2="10" y2="18" opacity="0.6" data-astro-cid-qlfjksao></line><line x1="26" y1="18" x2="34" y2="18" opacity="0.6" data-astro-cid-qlfjksao></line> <path d="M6 8 L6 6 L8 6" stroke-width="1.2" data-astro-cid-qlfjksao></path><path d="M28 6 L30 6 L30 8" stroke-width="1.2" data-astro-cid-qlfjksao></path> <path d="M30 28 L30 30 L28 30" stroke-width="1.2" data-astro-cid-qlfjksao></path><path d="M8 30 L6 30 L6 28" stroke-width="1.2" data-astro-cid-qlfjksao></path>      </svg> </div> <div class="min-w-0" data-astro-cid-qlfjksao> <span class="block text-[13px] font-semibold text-[#e4e4e7] tracking-wide" data-astro-cid-qlfjksao>EVALUATE</span> <span class="block text-[11px] text-[#71717a]" data-astro-cid-qlfjksao>Catch issues</span> </div> </div> <!-- Expand chevron --> <svg class="mobile-phase-chevron w-4 h-4 text-[#52525b] shrink-0 transition-transform duration-300" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <!-- Expandable links panel --> <div class="mobile-phase-links" style="max-height:0; overflow:hidden; transition: max-height 0.3s ease;" data-astro-cid-qlfjksao> <div class="pl-[50px] pr-1 pb-2 pt-1 space-y-0.5" data-astro-cid-qlfjksao> <a href="/platform/evaluate/error-feeds/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Error Feed</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/evaluate/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Evaluate</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/guard/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Protect</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> </div><div class="mobile-phase-card relative" data-phase-idx="3" data-astro-cid-qlfjksao> <!-- Phase header (tappable) --> <button class="mobile-phase-trigger w-full flex items-center gap-3 py-2 px-1 text-left group" data-astro-cid-qlfjksao> <!-- Phase node on timeline --> <div class="relative z-10 shrink-0 w-[38px] h-[38px] rounded-full bg-[#111111] border border-[#27272a] flex items-center justify-center transition-all duration-300 group-hover:border-[#52525b]" data-astro-cid-qlfjksao> <span class="text-[10px] font-mono font-bold text-[#71717a] transition-colors group-hover:text-[#a1a1aa]" data-astro-cid-qlfjksao>04</span> </div> <!-- Mini illustration + text --> <div class="flex-1 flex items-center gap-3 min-w-0" data-astro-cid-qlfjksao> <div class="shrink-0 w-9 h-9 flex items-center justify-center" data-astro-cid-qlfjksao> <svg class="w-8 h-8 text-[#52525b] transition-colors group-hover:text-[#71717a]" viewBox="0 0 36 36" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>         <path d="M4 24 Q4 8 18 3 Q32 8 32 24" stroke-width="1.2" data-astro-cid-qlfjksao></path> <path d="M8 24 Q8 11 18 7 Q28 11 28 24" stroke-dasharray="3 1.5" opacity="0.5" data-astro-cid-qlfjksao></path> <path d="M12 24 Q12 14 18 11 Q24 14 24 24" stroke-dasharray="1.5 1.5" opacity="0.3" data-astro-cid-qlfjksao></path> <path d="M18 15 L22 20 L22 27 L18 31 L14 27 L14 20 Z" stroke-width="1" opacity="0.7" data-astro-cid-qlfjksao></path> <circle cx="16" cy="29" r="1" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle> <circle cx="20" cy="29" r="1" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle> <circle cx="6" cy="12" r="1" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle><circle cx="30" cy="12" r="1" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle>    </svg> </div> <div class="min-w-0" data-astro-cid-qlfjksao> <span class="block text-[13px] font-semibold text-[#e4e4e7] tracking-wide" data-astro-cid-qlfjksao>OPTIMIZE</span> <span class="block text-[11px] text-[#71717a]" data-astro-cid-qlfjksao>Improve with data</span> </div> </div> <!-- Expand chevron --> <svg class="mobile-phase-chevron w-4 h-4 text-[#52525b] shrink-0 transition-transform duration-300" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <!-- Expandable links panel --> <div class="mobile-phase-links" style="max-height:0; overflow:hidden; transition: max-height 0.3s ease;" data-astro-cid-qlfjksao> <div class="pl-[50px] pr-1 pb-2 pt-1 space-y-0.5" data-astro-cid-qlfjksao> <a href="/platform/optimize/rl/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>AI Optimization</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> </div><div class="mobile-phase-card relative" data-phase-idx="4" data-astro-cid-qlfjksao> <!-- Phase header (tappable) --> <button class="mobile-phase-trigger w-full flex items-center gap-3 py-2 px-1 text-left group" data-astro-cid-qlfjksao> <!-- Phase node on timeline --> <div class="relative z-10 shrink-0 w-[38px] h-[38px] rounded-full bg-[#111111] border border-[#27272a] flex items-center justify-center transition-all duration-300 group-hover:border-[#52525b]" data-astro-cid-qlfjksao> <span class="text-[10px] font-mono font-bold text-[#71717a] transition-colors group-hover:text-[#a1a1aa]" data-astro-cid-qlfjksao>05</span> </div> <!-- Mini illustration + text --> <div class="flex-1 flex items-center gap-3 min-w-0" data-astro-cid-qlfjksao> <div class="shrink-0 w-9 h-9 flex items-center justify-center" data-astro-cid-qlfjksao> <svg class="w-8 h-8 text-[#52525b] transition-colors group-hover:text-[#71717a]" viewBox="0 0 36 36" fill="none" stroke="currentColor" stroke-width="1" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-qlfjksao>           <circle cx="18" cy="16" r="12" stroke-dasharray="3 1.5" opacity="0.4" data-astro-cid-qlfjksao></circle> <circle cx="18" cy="16" r="8" stroke-dasharray="2 1" opacity="0.6" data-astro-cid-qlfjksao></circle> <circle cx="18" cy="16" r="4" data-astro-cid-qlfjksao></circle> <line x1="18" y1="16" x2="18" y2="4" stroke-width="1.5" opacity="0.8" data-astro-cid-qlfjksao></line> <ellipse cx="18" cy="16" rx="12" ry="4" stroke-dasharray="2 1" opacity="0.3" data-astro-cid-qlfjksao></ellipse> <circle cx="24" cy="11" r="1.5" fill="currentColor" opacity="0.6" data-astro-cid-qlfjksao></circle><circle cx="12" cy="14" r="1" fill="currentColor" opacity="0.4" data-astro-cid-qlfjksao></circle><circle cx="26" cy="18" r="1" fill="currentColor" opacity="0.3" data-astro-cid-qlfjksao></circle> <ellipse cx="18" cy="32" rx="10" ry="2.5" stroke-dasharray="2 2" opacity="0.2" data-astro-cid-qlfjksao></ellipse> <path d="M8 32 L12 20" stroke-dasharray="1 2" opacity="0.15" data-astro-cid-qlfjksao></path><path d="M28 32 L24 20" stroke-dasharray="1 2" opacity="0.15" data-astro-cid-qlfjksao></path>  </svg> </div> <div class="min-w-0" data-astro-cid-qlfjksao> <span class="block text-[13px] font-semibold text-[#e4e4e7] tracking-wide" data-astro-cid-qlfjksao>MONITOR</span> <span class="block text-[11px] text-[#71717a]" data-astro-cid-qlfjksao>Insights in realtime</span> </div> </div> <!-- Expand chevron --> <svg class="mobile-phase-chevron w-4 h-4 text-[#52525b] shrink-0 transition-transform duration-300" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M19.5 8.25l-7.5 7.5-7.5-7.5" data-astro-cid-qlfjksao></path> </svg> </button> <!-- Expandable links panel --> <div class="mobile-phase-links" style="max-height:0; overflow:hidden; transition: max-height 0.3s ease;" data-astro-cid-qlfjksao> <div class="pl-[50px] pr-1 pb-2 pt-1 space-y-0.5" data-astro-cid-qlfjksao> <a href="/platform/monitor/tracing/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Tracing</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/monitor/dashboards/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Dashboards</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/monitor/alerting/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Alerting</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/platform/monitor/command-center/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[13px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Command Center</span> <svg class="w-3.5 h-3.5 text-[#52525b]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> </div> </div> </div> <!-- Audience --> <div class="mobile-nav-section" data-astro-cid-qlfjksao> <div class="text-[11px] text-[#71717a] font-medium tracking-wider uppercase mb-3" data-astro-cid-qlfjksao>Audience</div> <div class="space-y-1" data-astro-cid-qlfjksao> <a href="/enterprise/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span class="flex items-center gap-3" data-astro-cid-qlfjksao> <span class="w-8 h-8 rounded-lg bg-[#171717] border border-[#27272a] flex items-center justify-center" data-astro-cid-qlfjksao> <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M3.75 21h16.5M4.5 3h15M5.25 3v18m13.5-18v18M9 6.75h1.5m-1.5 3h1.5m-1.5 3h1.5m3-6H15m-1.5 3H15m-1.5 3H15M9 21v-3.375c0-.621.504-1.125 1.125-1.125h3.75c.621 0 1.125.504 1.125 1.125V21" data-astro-cid-qlfjksao></path> </svg> </span> <span class="flex flex-col" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Enterprise</span> <span class="text-[12px] text-[#71717a]" data-astro-cid-qlfjksao>Scale with confidence</span> </span> </span> <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/startups/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> <span class="flex items-center gap-3" data-astro-cid-qlfjksao> <span class="w-8 h-8 rounded-lg bg-[#171717] border border-[#27272a] flex items-center justify-center" data-astro-cid-qlfjksao> <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M15.59 14.37a6 6 0 01-5.84 7.38v-4.8m5.84-2.58a14.98 14.98 0 006.16-12.12A14.98 14.98 0 009.631 8.41m5.96 5.96a14.926 14.926 0 01-5.841 2.58m-.119-8.54a6 6 0 00-7.381 5.84h4.8m2.581-5.84a14.927 14.927 0 00-2.58 5.84m2.699 2.7c-.103.021-.207.041-.311.06a15.09 15.09 0 01-2.448-2.448 14.9 14.9 0 01.06-.312m-2.24 2.39a4.493 4.493 0 00-1.757 4.306 4.493 4.493 0 004.306-1.758M16.5 9a1.5 1.5 0 11-3 0 1.5 1.5 0 013 0z" data-astro-cid-qlfjksao></path> </svg> </span> <span class="flex flex-col" data-astro-cid-qlfjksao> <span data-astro-cid-qlfjksao>Startups</span> <span class="text-[12px] text-[#71717a]" data-astro-cid-qlfjksao>Move fast, stay safe</span> </span> </span> <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> <!-- Resources Sections --> <div class="mobile-nav-section" data-astro-cid-qlfjksao> <div class="text-[11px] text-[#71717a] font-medium tracking-wider uppercase mb-3" data-astro-cid-qlfjksao>LEARN</div> <div class="space-y-1" data-astro-cid-qlfjksao>  <a href="/customers/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Use Cases <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="/blog/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Blog <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="/ebooks/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> eBooks <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> <div class="ml-4 space-y-1 border-l border-[#27272a] pl-3" data-astro-cid-qlfjksao> <a href="/ebooks/mastering-ai-agent-evaluation/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> AI Agent Evaluation <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/ebooks/mastering-agentic-rag/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Agentic RAG Playbook <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="/ebooks/advanced-rag-patterns/" class="mobile-nav-link flex items-center justify-between px-3 py-2 text-[14px] text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Advanced RAG Patterns <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> <a href="/research/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Research <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="/changelog/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Changelog <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  </div> </div><div class="mobile-nav-section" data-astro-cid-qlfjksao> <div class="text-[11px] text-[#71717a] font-medium tracking-wider uppercase mb-3" data-astro-cid-qlfjksao>DEVELOPERS</div> <div class="space-y-1" data-astro-cid-qlfjksao>  <a href="https://docs.futureagi.com" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> Documentation <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="https://docs.futureagi.com/docs/api" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> API Reference <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="https://docs.futureagi.com/docs/sdk" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> SDK Reference <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="/integrations/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Integrations <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="/llm-cost-calculator/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> LLM Cost Calculator <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  <a href="/eval-tco-calculator/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Evaluation TCO Calculator <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a>  </div> </div> <!-- Featured --> <div class="mobile-nav-section" data-astro-cid-qlfjksao> <div class="text-[11px] text-[#71717a] font-medium tracking-wider uppercase mb-3" data-astro-cid-qlfjksao>Featured</div> <div class="space-y-3" data-astro-cid-qlfjksao> <a href="/ebooks/mastering-ai-agent-evaluation/" class="mobile-nav-link block px-3 py-3 bg-[#111111] hover:bg-[#171717] rounded-lg transition-colors border border-[#1f1f23]" data-astro-cid-qlfjksao> <div class="text-[14px] text-[#fafafa] font-medium mb-1" data-astro-cid-qlfjksao>Mastering AI Agent Evaluation</div> <p class="text-[12px] text-[#71717a] mb-2" data-astro-cid-qlfjksao>The complete guide to evaluating AI agents in production</p> <span class="inline-flex items-center gap-1 text-[12px] text-[#22c55e] font-medium" data-astro-cid-qlfjksao> <svg class="w-3 h-3" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M13.5 4.5L21 12m0 0l-7.5 7.5M21 12H3" data-astro-cid-qlfjksao></path> </svg> Download free </span> </a><a href="/ebooks/mastering-agentic-rag/" class="mobile-nav-link block px-3 py-3 bg-[#111111] hover:bg-[#171717] rounded-lg transition-colors border border-[#1f1f23]" data-astro-cid-qlfjksao> <div class="text-[14px] text-[#fafafa] font-medium mb-1" data-astro-cid-qlfjksao>The Agentic RAG Playbook</div> <p class="text-[12px] text-[#71717a] mb-2" data-astro-cid-qlfjksao>Enterprise RAG from theory to production-ready systems</p> <span class="inline-flex items-center gap-1 text-[12px] text-[#22c55e] font-medium" data-astro-cid-qlfjksao> <svg class="w-3 h-3" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M13.5 4.5L21 12m0 0l-7.5 7.5M21 12H3" data-astro-cid-qlfjksao></path> </svg> Download free </span> </a> </div> </div> <!-- Direct Links --> <div class="mobile-nav-section border-t border-[#27272a] pt-6" data-astro-cid-qlfjksao> <a href="/pricing/" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" data-astro-cid-qlfjksao> Pricing <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a><a href="https://docs.futureagi.com" class="mobile-nav-link flex items-center justify-between px-3 py-2.5 text-[15px] text-[#fafafa] hover:bg-[#171717] rounded-lg transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-qlfjksao> Docs <svg class="w-4 h-4 text-[#71717a]" fill="none" viewBox="0 0 24 24" stroke="currentColor" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" stroke-width="1.5" d="M8.25 4.5l7.5 7.5-7.5 7.5" data-astro-cid-qlfjksao></path> </svg> </a> </div> </div> <!-- Mobile CTAs --> <div class="flex flex-col gap-3 pt-4 border-t border-[#27272a]" data-astro-cid-qlfjksao> <span data-auth-slot class="contents" data-astro-cid-qlfjksao> <span data-auth-when="signed-out" class="contents" data-astro-cid-qlfjksao> <a href="https://app.futureagi.com/auth/jwt/register" target="_blank" rel="noopener noreferrer" class="inline-flex items-center justify-center font-medium rounded-full transition-all duration-200 focus:outline-none focus:ring-2 focus:ring-offset-2 focus:ring-offset-[#0a0a0a] bg-[#fafafa] text-[#0a0a0a] hover:bg-white focus:ring-white shadow-lg hover:shadow-xl text-[15px] px-6 py-3 gap-2.5 w-full"> 
Get Started - Free
 </a> </span> <span data-auth-when="signed-in" class="contents" data-astro-cid-qlfjksao> <a href="https://app.futureagi.com/dashboard/falcon-ai" target="_blank" rel="noopener noreferrer" class="inline-flex items-center justify-center font-medium rounded-full transition-all duration-200 focus:outline-none focus:ring-2 focus:ring-offset-2 focus:ring-offset-[#0a0a0a] bg-[#fafafa] text-[#0a0a0a] hover:bg-white focus:ring-white shadow-lg hover:shadow-xl text-[15px] px-6 py-3 gap-2.5 w-full"> 
Open dashboard
 </a> </span> </span> <a href="https://github.com/future-agi/future-agi" target="_blank" rel="noopener noreferrer" class="flex items-center justify-center gap-2 py-3 rounded-lg border border-[#27272a] bg-[#0f0f10] text-[14px] text-[#fafafa] hover:border-[#3f3f46] hover:bg-[#171717] transition-colors" data-astro-cid-qlfjksao> <svg class="w-4 h-4" fill="currentColor" viewBox="0 0 24 24" data-astro-cid-qlfjksao><path fill-rule="evenodd" clip-rule="evenodd" d="M12 2C6.477 2 2 6.484 2 12.017c0 4.425 2.865 8.18 6.839 9.504.5.092.682-.217.682-.483 0-.237-.008-.868-.013-1.703-2.782.605-3.369-1.343-3.369-1.343-.454-1.158-1.11-1.466-1.11-1.466-.908-.62.069-.608.069-.608 1.003.07 1.531 1.032 1.531 1.032.892 1.53 2.341 1.088 2.91.832.092-.647.35-1.088.636-1.338-2.22-.253-4.555-1.113-4.555-4.951 0-1.093.39-1.988 1.029-2.688-.103-.253-.446-1.272.098-2.65 0 0 .84-.27 2.75 1.026A9.564 9.564 0 0112 6.844c.85.004 1.705.115 2.504.337 1.909-1.296 2.747-1.027 2.747-1.027.546 1.379.202 2.398.1 2.651.64.7 1.028 1.595 1.028 2.688 0 3.848-2.339 4.695-4.566 4.943.359.309.678.92.678 1.855 0 1.338-.012 2.419-.012 2.747 0 .268.18.58.688.482A10.019 10.019 0 0022 12.017C22 6.484 17.522 2 12 2z" data-astro-cid-qlfjksao></path></svg> <span data-astro-cid-qlfjksao>Star on GitHub</span> <span class="inline-flex items-center gap-1 px-1.5 py-0.5 rounded text-[11px] font-medium tabular-nums bg-[#1a1a1e] text-[#e4e4e7]" data-astro-cid-qlfjksao> <svg class="w-3 h-3 text-[#facc15]" viewBox="0 0 24 24" fill="currentColor" data-astro-cid-qlfjksao><polygon points="12 2 15.09 8.26 22 9.27 17 14.14 18.18 21.02 12 17.77 5.82 21.02 7 14.14 2 9.27 8.91 8.26 12 2" data-astro-cid-qlfjksao></polygon></svg> <span data-stat="stars" data-stat-format="compact" data-astro-cid-qlfjksao>986</span> </span> </a> </div> </div> </div> </header> <!-- Search Modal (Command Palette Style) --> <div id="search-modal" class="fixed inset-0 z-[100] opacity-0 pointer-events-none transition-opacity duration-200" role="dialog" aria-modal="true" aria-labelledby="search-title" data-astro-cid-qlfjksao> <!-- Backdrop --> <div class="absolute inset-0 bg-black/60 backdrop-blur-sm overscroll-none" id="search-backdrop" data-astro-cid-qlfjksao></div> <!-- Modal Content --> <div class="relative flex items-start justify-center pt-[15vh] px-4" data-astro-cid-qlfjksao> <div class="w-full max-w-[600px] bg-[#0a0a0a] border border-[#27272a] rounded-xl shadow-2xl shadow-black/50 overflow-hidden transform scale-95 transition-transform duration-200" id="search-panel" data-astro-cid-qlfjksao> <!-- Search Input --> <div class="flex items-center gap-3 px-4 py-3 border-b border-[#27272a]" data-astro-cid-qlfjksao> <svg class="w-5 h-5 text-[#71717a] shrink-0" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="1.5" data-astro-cid-qlfjksao> <path stroke-linecap="round" stroke-linejoin="round" d="M21 21l-5.197-5.197m0 0A7.5 7.5 0 105.196 5.196a7.5 7.5 0 0010.607 10.607z" data-astro-cid-qlfjksao></path> </svg> <input id="search-input" type="text" placeholder="Search documentation, features, guides..." class="flex-1 bg-transparent text-[16px] text-[#fafafa] placeholder-[#52525b] outline-none" autocomplete="off" spellcheck="false" data-astro-cid-qlfjksao> <kbd class="px-2 py-1 text-[11px] font-medium text-[#71717a] bg-[#171717] border border-[#27272a] rounded" data-astro-cid-qlfjksao>
ESC
</kbd> </div> <!-- Search results (dynamically rendered). data-lenis-prevent opts wheel/touch events out of the global Lenis smooth-scroll engine so scrolling inside the dropdown doesn't scroll the page behind it. --> <div id="search-results" class="p-2 max-h-[400px] overflow-y-auto overscroll-contain" data-lenis-prevent data-astro-cid-qlfjksao></div> <!-- No results --> <div id="search-no-results" class="px-3 py-8 text-center hidden" data-astro-cid-qlfjksao> <p class="text-[13px] text-[#71717a]" data-astro-cid-qlfjksao>No results found</p> </div> <!-- Footer --> <div class="flex items-center justify-between px-4 py-3 border-t border-[#27272a] bg-[#0a0a0a]" data-astro-cid-qlfjksao> <div class="flex items-center gap-4 text-[12px] text-[#71717a]" data-astro-cid-qlfjksao> <span class="flex items-center gap-1" data-astro-cid-qlfjksao> <kbd class="px-1.5 py-0.5 bg-[#171717] border border-[#27272a] rounded text-[10px]" data-astro-cid-qlfjksao>↑</kbd> <kbd class="px-1.5 py-0.5 bg-[#171717] border border-[#27272a] rounded text-[10px]" data-astro-cid-qlfjksao>↓</kbd> <span class="ml-1" data-astro-cid-qlfjksao>Navigate</span> </span> <span class="flex items-center gap-1" data-astro-cid-qlfjksao> <kbd class="px-1.5 py-0.5 bg-[#171717] border border-[#27272a] rounded text-[10px]" data-astro-cid-qlfjksao>↵</kbd> <span class="ml-1" data-astro-cid-qlfjksao>Open</span> </span> </div> <span id="search-count" class="text-[12px] text-[#71717a]" data-astro-cid-qlfjksao></span> </div> <!-- Search index (injected at build time) --> <script id="search-index-data" type="application/json">[{"title":"Home","desc":"Future AGI - AI agent hallucination detection platform","href":"/","cat":"Pages"},{"title":"Pricing","desc":"Simple, transparent pricing. Start free, scale as you grow.","href":"/pricing/","cat":"Pages"},{"title":"Enterprise","desc":"Enterprise-grade AI safety at scale","href":"/enterprise/","cat":"Pages"},{"title":"Startups","desc":"$10K in free credits and 6 months Pro access","href":"/startups/","cat":"Pages"},{"title":"Roadmap","desc":"Public product roadmap - see what we're building next","href":"/roadmap/","cat":"Pages"},{"title":"Blog","desc":"Guides, engineering deep-dives, and product updates","href":"/blog/","cat":"Pages"},{"title":"Research","desc":"Papers on hallucination detection, evaluation, and guardrails","href":"/research/","cat":"Pages"},{"title":"Customers","desc":"Case studies from teams using Future AGI","href":"/customers/","cat":"Pages"},{"title":"eBooks","desc":"In-depth guides on AI agent evaluation and RAG","href":"/ebooks/","cat":"Pages"},{"title":"Guard","desc":"Block AI hallucinations in real-time with guardrails","href":"/platform/guard/","cat":"Platform"},{"title":"Evaluate","desc":"Run comprehensive evaluations with 20+ metrics","href":"/platform/evaluate/","cat":"Platform"},{"title":"Error Feed","desc":"Sentry-style error tracking for AI agents","href":"/platform/evaluate/error-feeds/","cat":"Platform"},{"title":"Simulations","desc":"Simulate thousands of multi-turn conversations","href":"/platform/simulate/","cat":"Platform"},{"title":"Scenarios","desc":"Define branching conversation test scenarios","href":"/platform/simulate/scenarios/","cat":"Platform"},{"title":"Synthetic Data","desc":"Generate diverse, realistic test data","href":"/platform/simulate/synthetic-data/","cat":"Platform"},{"title":"AI Optimization","desc":"Continuous improvement with reinforcement learning","href":"/platform/optimize/rl/","cat":"Platform"},{"title":"Tracing","desc":"End-to-end request tracing for AI agents","href":"/platform/monitor/tracing/","cat":"Platform"},{"title":"Dashboards","desc":"Custom dashboards with drag-and-drop widgets","href":"/platform/monitor/dashboards/","cat":"Platform"},{"title":"Alerting","desc":"AI-powered alerts for anomalies and hallucination spikes","href":"/platform/monitor/alerting/","cat":"Platform"},{"title":"Guardrails (Monitor)","desc":"Real-time guardrail monitoring and block rate insights","href":"/platform/monitor/guardrails/","cat":"Platform"},{"title":"Datasets","desc":"Manage and version evaluation datasets","href":"/platform/agents/datasets/","cat":"Platform"},{"title":"Experiments","desc":"Structured experiments across models and prompts","href":"/platform/agents/experiments/","cat":"Platform"},{"title":"Agent IDE","desc":"Build & test AI agents visually","href":"/platform/agents/ide/","cat":"Platform"},{"title":"Introduction","desc":"The complete platform to test, guard, and monitor AI agents. Build self-improving agents that ship smarter with every version.","href":"https://docs.futureagi.com/docs","cat":"Docs"},{"title":"Overview","desc":"Deploy the full Future AGI platform on your own infrastructure using Docker Compose. Follow the step-by-step guide to get all services running locally.","href":"https://docs.futureagi.com/docs/self-hosting","cat":"Docs"},{"title":"Requirements","desc":"Hardware sizing tiers, supported platforms, OS compatibility, and network port requirements before deploying Future AGI with Docker Compose.","href":"https://docs.futureagi.com/docs/self-hosting/requirements","cat":"Docs"},{"title":"Docker Compose","desc":"Deploy the full Future AGI stack with Docker Compose — all 21 services, dev overlay with hot reload, and frontend-only mode pointing at a remote backend.","href":"https://docs.futureagi.com/docs/self-hosting/docker-compose","cat":"Docs"},{"title":"Environment Variables","desc":"Full .env reference for self-hosted Future AGI — secrets, database credentials, runtime flags, LLM provider keys, email, and frontend build-time configuration.","href":"https://docs.futureagi.com/docs/self-hosting/environment","cat":"Docs"},{"title":"System Configuration","desc":"Configure the LLM gateway config.yaml with provider API keys, set up PeerDB Postgres-to-ClickHouse CDC mirrors, and tune Temporal worker concurrency.","href":"https://docs.futureagi.com/docs/self-hosting/configuration","cat":"Docs"},{"title":"User Management","desc":"Create user accounts, reset passwords, and manage roles in self-hosted Future AGI — via Mailgun email flow or directly through the Django admin shell.","href":"https://docs.futureagi.com/docs/self-hosting/user-management","cat":"Docs"},{"title":"Production","desc":"Production readiness checklist — replace secrets, configure TLS, set up managed data stores, run Postgres/ClickHouse/MinIO backups, and follow the upgrade runbook.","href":"https://docs.futureagi.com/docs/self-hosting/production","cat":"Docs"},{"title":"Troubleshooting","desc":"Debug self-hosted Future AGI — symptoms, causes, and fixes for startup failures, network issues, PeerDB CDC errors, Temporal worker problems, and post-upgrade breaks.","href":"https://docs.futureagi.com/docs/self-hosting/troubleshooting","cat":"Docs"},{"title":"Quickstart","desc":"The complete platform to test, guard, and monitor AI agents. Build self-improving agents that ship smarter with every version.","href":"https://docs.futureagi.com/docs","cat":"Docs"},{"title":"Create Prompts","desc":"Create and manage AI prompts in Future AGI's Prompt Workbench. Design, test, version, and optimize prompts with built-in model selection and evaluation.","href":"https://docs.futureagi.com/docs/quickstart/prompts","cat":"Docs"},{"title":"Generate Synthetic Data","desc":"Generate synthetic datasets with Future AGI. Define schemas, column types, and constraints to create realistic data for training and evaluation.","href":"https://docs.futureagi.com/docs/quickstart/generate-synthetic-data","cat":"Docs"},{"title":"Running Evals in Simulation","desc":"Run evaluations in Future AGI simulations to test AI agents against simulated customers and score interactions for quality and context retention.","href":"https://docs.futureagi.com/docs/quickstart/running-evals-in-simulation","cat":"Docs"},{"title":"Agent Command Center","desc":"Make your first LLM request through the Future AGI Agent Command Center in under 5 minutes with no framework changes required.","href":"https://docs.futureagi.com/docs/quickstart/command-center-gateway","cat":"Docs"},{"title":"Setup Observability","desc":"Set up Future AGI Observe for production monitoring. Configure auto-instrumented tracing for OpenAI, Anthropic, LangChain, and other LLM frameworks.","href":"https://docs.futureagi.com/docs/quickstart/setup-observability","cat":"Docs"},{"title":"Annotations","desc":"Get started with Future AGI annotations in 5 minutes: create annotation labels, set up a queue, add items, and start annotating.","href":"https://docs.futureagi.com/docs/quickstart/annotations","cat":"Docs"},{"title":"Setup MCP Server","desc":"Set up the Future AGI MCP Server to interact with the platform via natural language from Claude, Cursor, or VS Code using Model Context Protocol.","href":"https://docs.futureagi.com/docs/quickstart/setup-mcp-server","cat":"Docs"},{"title":"Release Notes","desc":"Latest Future AGI release notes covering new features, improvements, and bug fixes across datasets, evaluations, simulation, and observability products.","href":"https://docs.futureagi.com/docs/release-notes","cat":"Docs"},{"title":"Overview","desc":"Design and run multi-step AI agent workflows on a drag-and-drop canvas. Chain LLM calls, embed sub-agents, version changes, and trace every execution.","href":"https://docs.futureagi.com/docs/agent-playground","cat":"Docs"},{"title":"Understanding Agent Playground","desc":"Learn the core building blocks of Agent Playground: graphs, LLM nodes, subgraph nodes, ports, edges, and node templates.","href":"https://docs.futureagi.com/docs/agent-playground/concepts/understanding-agent-playground","cat":"Docs"},{"title":"Versions & Execution","desc":"Understand Agent Playground's draft/non-draft version lifecycle, topological execution model, node states, and data routing between nodes.","href":"https://docs.futureagi.com/docs/agent-playground/concepts/versions-and-execution","cat":"Docs"},{"title":"Create a Graph","desc":"Create a new agent graph in Agent Playground, set metadata, manage draft and saved versions, and roll back to previous workflow snapshots.","href":"https://docs.futureagi.com/docs/agent-playground/features/create-graph","cat":"Docs"},{"title":"Build a Workflow","desc":"Add LLM Prompt and Agent nodes to the canvas, configure models and parameters, draw edges, and set global variables in Agent Playground.","href":"https://docs.futureagi.com/docs/agent-playground/features/build-workflow","cat":"Docs"},{"title":"Run & Monitor","desc":"Execute AI agent workflows, watch per-node status in real time, inspect input/output data for each step, and browse full execution history.","href":"https://docs.futureagi.com/docs/agent-playground/features/run-and-monitor","cat":"Docs"},{"title":"Overview","desc":"Capture human feedback on AI outputs using labels, queues, and scores across traces, spans, sessions, datasets, prototypes, and simulations.","href":"https://docs.futureagi.com/docs/annotations","cat":"Docs"},{"title":"Scores","desc":"Understand the Score model — the unified annotation primitive storing label values, annotator, source type, and queue context across all entity types.","href":"https://docs.futureagi.com/docs/annotations/concepts/scores","cat":"Docs"},{"title":"Labels","desc":"Create and configure annotation labels: categorical, numeric, text, star rating, and thumbs up/down. Reusable across all queues in your organization.","href":"https://docs.futureagi.com/docs/annotations/features/labels","cat":"Docs"},{"title":"Queues","desc":"Create annotation queues with round-robin or load-balanced assignment, multi-annotator support, reservation timeouts, and review workflows.","href":"https://docs.futureagi.com/docs/annotations/features/queues","cat":"Docs"},{"title":"Add Items to Queues","desc":"Add traces, spans, sessions, dataset rows, prototype runs, and simulation calls to annotation queues for structured human review in Future AGI.","href":"https://docs.futureagi.com/docs/annotations/features/add-items","cat":"Docs"},{"title":"Annotate Items","desc":"Use the annotation workspace to label traces, sessions, and datasets with categorical, numeric, star, and thumbs inputs plus keyboard shortcuts.","href":"https://docs.futureagi.com/docs/annotations/features/annotate","cat":"Docs"},{"title":"Inline Annotations","desc":"Score traces, sessions, dataset rows, and prototype runs directly from their detail views without setting up an annotation queue.","href":"https://docs.futureagi.com/docs/annotations/features/inline","cat":"Docs"},{"title":"Analytics & Agreement","desc":"Track queue progress, annotator throughput, label distribution, and inter-annotator agreement using Cohen's and Fleiss' Kappa metrics.","href":"https://docs.futureagi.com/docs/annotations/features/analytics","cat":"Docs"},{"title":"Export Annotations","desc":"Export completed annotation queue results to a Future AGI dataset or download as JSON/CSV for fine-tuning, evaluation, and offline analysis.","href":"https://docs.futureagi.com/docs/annotations/features/export","cat":"Docs"},{"title":"Automation Rules","desc":"Create condition-based rules to automatically add matching traces, spans, or sessions to annotation queues without manual curation.","href":"https://docs.futureagi.com/docs/annotations/features/automation","cat":"Docs"},{"title":"Python SDK","desc":"Log annotations via DataFrame, retrieve labels, list projects, and submit human feedback to traces using the FutureAGI Python SDK.","href":"https://docs.futureagi.com/docs/annotations/sdk/python","cat":"Docs"},{"title":"JavaScript SDK","desc":"Log annotations, manage queues, submit scores, and export results using the FutureAGI JavaScript/TypeScript SDK's Annotation and AnnotationQueue classes.","href":"https://docs.futureagi.com/docs/annotations/sdk/javascript","cat":"Docs"},{"title":"Annotation Queue Using SDK","desc":"Create queues, manage labels, add items, submit annotations, track progress, and export results programmatically with the Future AGI Python SDK.","href":"https://docs.futureagi.com/docs/annotations/sdk/annotation-queue-using-sdk","cat":"Docs"},{"title":"Overview","desc":"A unified API gateway for 100+ LLM providers with built-in guardrails, intelligent routing, caching, cost controls, and full observability.","href":"https://docs.futureagi.com/docs/command-center","cat":"Docs"},{"title":"How it works","desc":"How Agent Command Center works — request pipeline, plugin execution order, virtual keys, per-org multi-tenancy, caching, guardrails, and cost tracking.","href":"https://docs.futureagi.com/docs/command-center/concepts/core","cat":"Docs"},{"title":"Virtual keys & access control","desc":"Manage sk-agentcc-* virtual keys in Agent Command Center — restrict models, providers, and IPs with RBAC roles and enforce per-key rate limits and budgets.","href":"https://docs.futureagi.com/docs/command-center/concepts/virtual-keys","cat":"Docs"},{"title":"Configuration","desc":"Learn how Agent Command Center configuration works — hierarchy from request headers down to global defaults, SDK config objects, and real-time org-level settings.","href":"https://docs.futureagi.com/docs/command-center/concepts/configuration","cat":"Docs"},{"title":"Platform integration","desc":"How Agent Command Center feeds signals into Future AGI Observe, Evaluate, Protect, and Experiment to close the loop between production traffic and model quality.","href":"https://docs.futureagi.com/docs/command-center/concepts/platform-integration","cat":"Docs"},{"title":"Supported providers","desc":"Connect to 20+ cloud and self-hosted LLM providers via a unified OpenAI-compatible API. Add a provider once and switch by changing the model name.","href":"https://docs.futureagi.com/docs/command-center/features/providers","cat":"Docs"},{"title":"Self-hosted models","desc":"Route requests to Ollama, vLLM, LM Studio, or any OpenAI-compatible inference server alongside cloud providers. All gateway features apply to local models.","href":"https://docs.futureagi.com/docs/command-center/features/self-hosted-models","cat":"Docs"},{"title":"Endpoints overview","desc":"Full list of 97 endpoints across 20+ categories in Agent Command Center — inference under /v1/ and admin endpoints under /-/ with OpenAI-compatible format.","href":"https://docs.futureagi.com/docs/command-center/api/endpoints","cat":"Docs"},{"title":"Chat completions","desc":"POST /v1/chat/completions through Agent Command Center — OpenAI-compatible with routing, caching, guardrails, streaming, function calling, and vision.","href":"https://docs.futureagi.com/docs/command-center/api/chat","cat":"Docs"},{"title":"Embeddings & reranking","desc":"Generate text embeddings and rerank documents via Agent Command Center. OpenAI-compatible format with caching, cost tracking, failover, and rate limiting.","href":"https://docs.futureagi.com/docs/command-center/api/embeddings","cat":"Docs"},{"title":"Media endpoints","desc":"Text-to-speech, speech-to-text, audio translation, and image generation via Agent Command Center with caching, cost tracking, and provider failover.","href":"https://docs.futureagi.com/docs/command-center/api/media","cat":"Docs"},{"title":"Assistants API","desc":"Proxy the OpenAI Assistants API through Agent Command Center with routing, cost tracking, and logging for threads, runs, and tool use.","href":"https://docs.futureagi.com/docs/command-center/api/assistants","cat":"Docs"},{"title":"Files & vector stores","desc":"Upload files and manage vector stores via Agent Command Center. Supports Assistants API file search, fine-tuning datasets, and batch processing workflows.","href":"https://docs.futureagi.com/docs/command-center/api/files","cat":"Docs"},{"title":"Async & batch","desc":"Run async inference jobs and bulk batch requests through Agent Command Center. Poll job IDs for results or submit large batches at reduced cost.","href":"https://docs.futureagi.com/docs/command-center/api/async-batch","cat":"Docs"},{"title":"Request & response headers","desc":"Reference for all x-agentcc-* headers in Agent Command Center — control caching, sessions, and routing per request; read cost, latency, and provider in responses.","href":"https://docs.futureagi.com/docs/command-center/api/headers","cat":"Docs"},{"title":"Routing & reliability","desc":"Distribute LLM traffic across providers with automatic failover, retries, circuit breaking, and weight-based routing for high availability.","href":"https://docs.futureagi.com/docs/command-center/features/routing","cat":"Docs"},{"title":"Guardrails","desc":"Enforce PII detection, prompt injection blocking, and content moderation on LLM traffic. 18+ built-in guardrail types with enforce, monitor, or log modes.","href":"https://docs.futureagi.com/docs/command-center/features/guardrails","cat":"Docs"},{"title":"Caching","desc":"Cache LLM responses at the gateway with exact match and semantic caching. Cut provider costs and latency for repeated queries across all providers.","href":"https://docs.futureagi.com/docs/command-center/features/caching","cat":"Docs"},{"title":"Rate limiting","desc":"Set per-key, per-org, and global RPM limits. Enforce monthly spend budgets and per-key credit balances to prevent runaway costs and protect provider quotas.","href":"https://docs.futureagi.com/docs/command-center/features/rate-limiting","cat":"Docs"},{"title":"Cost tracking","desc":"Automatically track LLM spend per request via x-agentcc-cost header. Attribute costs by team, feature, or user and configure budget threshold alerts.","href":"https://docs.futureagi.com/docs/command-center/features/cost-tracking","cat":"Docs"},{"title":"Observability","desc":"Log every LLM request automatically with token counts, cost, latency, cache status, and guardrail results. Export metrics to Prometheus and OpenTelemetry.","href":"https://docs.futureagi.com/docs/command-center/features/observability","cat":"Docs"},{"title":"Shadow experiments","desc":"Silently mirror production LLM requests to a shadow model without affecting users. Compare cost, latency, and quality on real traffic before switching.","href":"https://docs.futureagi.com/docs/command-center/features/shadow-experiments","cat":"Docs"},{"title":"Webhooks","desc":"Receive real-time HTTP notifications for completed requests, guardrail triggers, budget alerts, and errors. Build integrations and audit pipelines.","href":"https://docs.futureagi.com/docs/command-center/features/webhooks","cat":"Docs"},{"title":"Custom Properties","desc":"Define typed metadata schemas (String, Number, Boolean, Enum) to annotate gateway requests. Filter and segment logs by team, environment, or cost center.","href":"https://docs.futureagi.com/docs/command-center/features/custom-properties","cat":"Docs"},{"title":"MCP & A2A","desc":"Connect AI agents via MCP (Model Context Protocol) and A2A (Agent-to-Agent) to aggregate tools, delegate tasks, and build multi-agent networks.","href":"https://docs.futureagi.com/docs/command-center/features/mcp-a2a","cat":"Docs"},{"title":"Organization management","desc":"Create and manage isolated organizations in Agent Command Center — configure providers, routing rules, rate limits, budgets, and API keys per org.","href":"https://docs.futureagi.com/docs/command-center/admin/organizations","cat":"Docs"},{"title":"Self-hosted","desc":"Deploy Agent Command Center on your own infrastructure via Docker or Go binary — full control over data residency, routing, failover, caching, and rate limiting.","href":"https://docs.futureagi.com/docs/command-center/deployment/self-hosted","cat":"Docs"},{"title":"Error handling","desc":"Reference for Agent Command Center error JSON format, 4xx and 5xx status codes with machine-readable codes, and retry strategies.","href":"https://docs.futureagi.com/docs/command-center/guides/errors","cat":"Docs"},{"title":"Troubleshooting","desc":"Debug checklist and fixes for common Agent Command Center issues: model not found, 429 rate limits, provider 502 errors, and slow responses.","href":"https://docs.futureagi.com/docs/command-center/guides/troubleshooting","cat":"Docs"},{"title":"Overview","desc":"Structured tables of examples for prompts, evaluations, and experiments. Create from file uploads, SDK, production traces, or synthetic generation.","href":"https://docs.futureagi.com/docs/dataset","cat":"Docs"},{"title":"Understanding Datasets","desc":"Each row is one example; each column is an attribute. Datasets are the foundation for running prompts, evals, experiments, and optimizations in Future AGI.","href":"https://docs.futureagi.com/docs/dataset/concept/understanding-dataset","cat":"Docs"},{"title":"Static Columns","desc":"Dataset columns for storing fixed test inputs, expected outputs, labels, and metadata. Supports 9 data types including text, JSON, image, and audio.","href":"https://docs.futureagi.com/docs/dataset/concept/static-column","cat":"Docs"},{"title":"Dynamic Columns","desc":"Dataset columns auto-generated by running LLM prompts, evaluations, vector retrieval, entity extraction, or custom Python code against every row.","href":"https://docs.futureagi.com/docs/dataset/concept/dynamic-column","cat":"Docs"},{"title":"Synthetic Data","desc":"Generate schema-driven test datasets without using real user data. Define column types, constraints, and descriptions, then generate rows using Future AGI.","href":"https://docs.futureagi.com/docs/dataset/concept/synthetic-data","cat":"Docs"},{"title":"Create New Dataset","desc":"Create a dataset from CSV, Hugging Face, production traces, or synthetic generation. Use it as the container for prompts, evals, and experiments.","href":"https://docs.futureagi.com/docs/dataset/features/create","cat":"Docs"},{"title":"Add Rows to Dataset","desc":"Add data points to an existing dataset manually, from another dataset, Hugging Face, from production traces, or by generating synthetic rows.","href":"https://docs.futureagi.com/docs/dataset/features/add-rows","cat":"Docs"},{"title":"Add Columns to Dataset","desc":"Add static columns for fixed values or dynamic columns whose values are computed from other columns or external operations.","href":"https://docs.futureagi.com/docs/dataset/features/add-columns","cat":"Docs"},{"title":"Run Prompt in Dataset","desc":"Add a dynamic column to your dataset by running an LLM, TTS, STT, or image generation model on every row using a prompt with column placeholders.","href":"https://docs.futureagi.com/docs/dataset/features/run-prompt","cat":"Docs"},{"title":"Experiments in Dataset","desc":"Test different prompt and model combinations on the same dataset. Score outputs with built-in evals and compare results side by side.","href":"https://docs.futureagi.com/docs/dataset/features/experiments","cat":"Docs"},{"title":"Add Annotation","desc":"Annotations are essential for refining datasets, evaluating model outputs, and improving the quality of AI-generated responses.","href":"https://docs.futureagi.com/docs/dataset/features/annotate","cat":"Docs"},{"title":"Overview","desc":"Automatically detect, cluster, score, and triage errors in your AI agent traces, without any configuration beyond standard tracing.","href":"https://docs.futureagi.com/docs/error-feed","cat":"Docs"},{"title":"How It Works","desc":"The mental model behind Error Feed: how traces become analyzed issues, how similar errors are grouped, and how findings surface in the UI.","href":"https://docs.futureagi.com/docs/error-feed/concepts/how-it-works","cat":"Docs"},{"title":"Error Taxonomy","desc":"Reference for the five categories of errors Error Feed detects in AI agent traces, with every subcategory and error type defined.","href":"https://docs.futureagi.com/docs/error-feed/concepts/taxonomy","cat":"Docs"},{"title":"Scoring","desc":"The four quality metrics Error Feed uses to score every analyzed trace, what each one measures, how scores are assigned, and how to interpret them.","href":"https://docs.futureagi.com/docs/error-feed/concepts/scoring","cat":"Docs"},{"title":"Severity and Status","desc":"Error Feed severity levels classify how critical each issue is, and status labels track issues through triage from Open to Resolved.","href":"https://docs.futureagi.com/docs/error-feed/concepts/severity-and-status","cat":"Docs"},{"title":"The Feed","desc":"How to read the Error Feed issue list: filters, the stats bar, table columns, trend sparklines, and time range controls.","href":"https://docs.futureagi.com/docs/error-feed/features/the-feed","cat":"Docs"},{"title":"Issue Overview","desc":"A walkthrough of the Overview tab on an issue detail page: every section explained, from the header to the AI-generated recommendations.","href":"https://docs.futureagi.com/docs/error-feed/features/issue-overview","cat":"Docs"},{"title":"Traces","desc":"How to use the Traces tab on an issue detail page to navigate every trace in a cluster and understand the distribution of failures.","href":"https://docs.futureagi.com/docs/error-feed/features/traces","cat":"Docs"},{"title":"State Graph","desc":"How to read the State Graph tab, the agent decision flow diagram showing where traces diverge between success and failure paths.","href":"https://docs.futureagi.com/docs/error-feed/features/state-graph","cat":"Docs"},{"title":"Trends","desc":"How to use the Trends tab (Events Over Time, Score Trends, and the Activity Heatmap) to understand how an issue is evolving.","href":"https://docs.futureagi.com/docs/error-feed/features/trends","cat":"Docs"},{"title":"Metadata Panel","desc":"The right-side metadata panel on an Error Feed issue page covers triage controls, cluster stats, timeline, evaluations, and co-occurring issues.","href":"https://docs.futureagi.com/docs/error-feed/features/metadata-panel","cat":"Docs"},{"title":"Triage Workflow","desc":"How to move issues through the Error Feed triage workflow: resolving, acknowledging, ignoring, assigning, and escalating.","href":"https://docs.futureagi.com/docs/error-feed/features/triage-workflow","cat":"Docs"},{"title":"Deep Analysis","desc":"On-demand investigation that runs deeper root cause analysis on an Error Feed issue's trace and produces richer findings than the continuous scan.","href":"https://docs.futureagi.com/docs/error-feed/features/deep-analysis","cat":"Docs"},{"title":"Linear Integration","desc":"Create Linear tickets directly from Error Feed issues to link your AI error monitoring to your engineering workflow and track fixes.","href":"https://docs.futureagi.com/docs/error-feed/features/linear-integration","cat":"Docs"},{"title":"Sampling","desc":"How sampling rate controls what percentage of traces Error Feed analyzes, and how to configure it per project in Observe settings.","href":"https://docs.futureagi.com/docs/error-feed/features/sampling","cat":"Docs"},{"title":"Overview","desc":"Measure and compare the quality of prompts and agents across datasets, simulations, and experiments using built-in or custom eval templates.","href":"https://docs.futureagi.com/docs/evaluation","cat":"Docs"},{"title":"Understanding Evaluation","desc":"Covers how evaluation works in Future AGI: templates, judge models, eval types, results, and where evaluations run in the platform and SDK.","href":"https://docs.futureagi.com/docs/evaluation/concepts/understanding-evaluation","cat":"Docs"},{"title":"Eval Types","desc":"The four evaluation methods in Future AGI: LLM as Judge, Deterministic, Statistical Metric, and LLM as Ranker, and how modality affects which ones apply.","href":"https://docs.futureagi.com/docs/evaluation/concepts/eval-types","cat":"Docs"},{"title":"Eval Templates","desc":"Explains what evaluation templates are, the difference between built-in and custom templates, and how output types determine what an eval returns.","href":"https://docs.futureagi.com/docs/evaluation/concepts/eval-templates","cat":"Docs"},{"title":"Judge Models","desc":"Explains what a judge model is, how it scores AI responses, and how to choose the right judge model for your specific evaluation use case.","href":"https://docs.futureagi.com/docs/evaluation/concepts/judge-models","cat":"Docs"},{"title":"Eval Results","desc":"Understand what evaluation results contain, how to read them, and how results are stored and aggregated across runs in Future AGI.","href":"https://docs.futureagi.com/docs/evaluation/concepts/eval-results","cat":"Docs"},{"title":"Built-in Evals","desc":"Complete reference for all built-in evaluation templates available on the Future AGI platform, with quick access to metrics by name.","href":"https://docs.futureagi.com/docs/evaluation/builtin","cat":"Docs"},{"title":"Evaluate via Platform & SDK","desc":"Run evaluations using the Future AGI platform UI or the Python SDK. Choose individual templates or batch runs for scalable model assessment.","href":"https://docs.futureagi.com/docs/evaluation/features/evaluate","cat":"Docs"},{"title":"Create Custom Evals","desc":"Define custom evaluation criteria and rules tailored to your use case, extending beyond the built-in templates available in Future AGI.","href":"https://docs.futureagi.com/docs/evaluation/features/custom","cat":"Docs"},{"title":"Use Custom Models","desc":"Use your own or third-party models for evaluations in Future AGI via supported providers or a custom API endpoint with full configuration control.","href":"https://docs.futureagi.com/docs/evaluation/features/custom-models","cat":"Docs"},{"title":"Future AGI Models","desc":"Future AGI's proprietary judge models are trained on diverse datasets to perform accurate evaluations and score AI outputs.","href":"https://docs.futureagi.com/docs/evaluation/features/futureagi-models","cat":"Docs"},{"title":"Evaluate CI/CD Pipeline","desc":"Run Future AGI evaluations in your CI/CD pipeline to assess model performance on every pull request and keep quality checks consistent before deployment.","href":"https://docs.futureagi.com/docs/evaluation/features/cicd","cat":"Docs"},{"title":"Overview","desc":"An AI copilot embedded in the Future AGI dashboard that handles platform tasks, runs analysis, and answers questions through natural language.","href":"https://docs.futureagi.com/docs/falcon-ai","cat":"Docs"},{"title":"Using Falcon AI","desc":"Open Falcon AI from any page, ask questions, upload files, and get streaming responses with tool calls and completion cards.","href":"https://docs.futureagi.com/docs/falcon-ai/features/chat","cat":"Docs"},{"title":"Skill Builder","desc":"Use built-in skills for common workflows or create custom slash commands that package multi-step instructions for your team.","href":"https://docs.futureagi.com/docs/falcon-ai/features/skills","cat":"Docs"},{"title":"MCP Connectors","desc":"Connect external MCP servers to Falcon AI to use tools from services like Linear, Slack, GitHub, Sentry, and custom APIs.","href":"https://docs.futureagi.com/docs/falcon-ai/features/mcp-connectors","cat":"Docs"},{"title":"Overview","desc":"Store your organization’s content in Future AGI to ground synthetic data generation and evaluations in real source material.","href":"https://docs.futureagi.com/docs/knowledge-base","cat":"Docs"},{"title":"Understanding Knowledge Base","desc":"Explains what a Knowledge Base is, what content types are supported, and how files are indexed and processed in Future AGI.","href":"https://docs.futureagi.com/docs/knowledge-base/concepts/concept","cat":"Docs"},{"title":"Create KB Using SDK","desc":"Create and manage Knowledge Bases programmatically with the Future AGI Python SDK: create, update, add or remove files, and delete KBs.","href":"https://docs.futureagi.com/docs/knowledge-base/features/sdk","cat":"Docs"},{"title":"Create KB Using UI","desc":"Create and populate a Knowledge Base from the Future AGI platform: name it, upload documents, and wait for processing to finish.","href":"https://docs.futureagi.com/docs/knowledge-base/features/ui","cat":"Docs"},{"title":"Overview","desc":"Monitor and evaluate LLM applications in production with real-time tracing, session analysis, cost tracking, and alerting.","href":"https://docs.futureagi.com/docs/observe","cat":"Docs"},{"title":"Understanding Observability","desc":"Core concepts behind LLM observability in Future AGI: what gets captured, how data is structured, and why monitoring matters for AI applications.","href":"https://docs.futureagi.com/docs/tracing/concepts","cat":"Docs"},{"title":"What are Traces?","desc":"A trace in Future AGI represents the full execution flow of an AI request, composed of spans that capture each LLM call, tool use, and retrieval step.","href":"https://docs.futureagi.com/docs/tracing/concepts/traces","cat":"Docs"},{"title":"What are Spans?","desc":"Understand spans in Future AGI tracing. Learn about span types including LLM, tool, chain, retriever, and embedding spans with their attributes.","href":"https://docs.futureagi.com/docs/tracing/concepts/spans","cat":"Docs"},{"title":"What is OpenTelemetry?","desc":"Learn how Future AGI uses OpenTelemetry for vendor-neutral, high-performance tracing of AI applications with standardized telemetry collection.","href":"https://docs.futureagi.com/docs/tracing/concepts/otel","cat":"Docs"},{"title":"What is traceAI?","desc":"Learn about traceAI, Future AGI's open-source package for standardized AI application tracing built on OpenTelemetry with framework-specific instrumentors.","href":"https://docs.futureagi.com/docs/tracing/concepts/traceai","cat":"Docs"},{"title":"Set Up Observability","desc":"Instrument your application and send traces to an Observe project so you can monitor LLM calls, latency, and cost in one place.","href":"https://docs.futureagi.com/docs/observe/features/quickstart","cat":"Docs"},{"title":"Run Evals on Traces","desc":"Run automated quality checks on traced spans in Observe: filter spans, choose historic or continuous runs, set sampling, and attach preset or custom evals.","href":"https://docs.futureagi.com/docs/observe/features/evals","cat":"Docs"},{"title":"Sessions","desc":"Group traces into sessions so you can view and analyze multi-turn conversations, chatbot flows, and per-session metrics in Observe.","href":"https://docs.futureagi.com/docs/observe/features/session","cat":"Docs"},{"title":"Users","desc":"View all traces, sessions, and metrics per end user in one place so you can debug, analyze behavior, and optimize at the user level.","href":"https://docs.futureagi.com/docs/observe/features/users","cat":"Docs"},{"title":"Alerts & Monitors","desc":"Define monitors on Observe project metrics (system or evaluation) and get notified by email or Slack when values cross a threshold.","href":"https://docs.futureagi.com/docs/observe/features/alerts","cat":"Docs"},{"title":"Voice Observability","desc":"Connect a voice provider like Vapi or Retell and get call logs as traces in Observe without any SDK instrumentation or code changes.","href":"https://docs.futureagi.com/docs/observe/features/voice","cat":"Docs"},{"title":"Dashboards","desc":"Build custom dashboards with widgets to visualize your Observe project metrics, traces, and performance data in one place.","href":"https://docs.futureagi.com/docs/observe/features/dashboard","cat":"Docs"},{"title":"Set Up Tracing","desc":"Connect your application to Future AGI by registering a tracer provider and adding instrumentation with auto-instrumentors or manual OpenTelemetry spans.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/set-up-tracing","cat":"Docs"},{"title":"Instrument with traceAI Helpers","desc":"Future AGI's traceAI library offers convenient abstractions that streamline your manual instrumentation process for LLM and agent tracing.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/instrument-with-traceai-helpers","cat":"Docs"},{"title":"Get Current Tracer and Span","desc":"Access the active span or tracer at any point in your code to enrich traces with additional attributes and context in Future AGI.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/get-current-span-context","cat":"Docs"},{"title":"Enriching Spans with Attributes, Metadata, and Tags","desc":"Enrich spans with custom attributes, metadata, tags, session IDs, user IDs, and prompt templates beyond what standard auto-instrumentation captures.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/add-attributes-metadata-tags","cat":"Docs"},{"title":"Logging Prompt Templates & Variables","desc":"Attach prompt template data to spans so Future AGI can surface it in the prompt playground for testing changes without deploying.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/log-prompt-templates","cat":"Docs"},{"title":"Events, Exceptions, and Status","desc":"Add OpenTelemetry events, exceptions, and status codes to spans in Future AGI to capture structured lifecycle information and error diagnostics.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/add-events-exceptions-status","cat":"Docs"},{"title":"Set Session ID and User ID","desc":"Add SessionID and UserID as span attributes to group and filter traces by conversation session and end user in Future AGI.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/set-session-user-id","cat":"Docs"},{"title":"Tool Spans Creation","desc":"Manually trace tool functions alongside LLM calls by creating spans that capture inputs, outputs, and key events in Future AGI.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/create-tool-spans","cat":"Docs"},{"title":"Mask Span Attributes","desc":"Redact sensitive inputs, outputs, images, and embeddings from spans before export, using environment variables or TraceConfig in code.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/mask-span-attributes","cat":"Docs"},{"title":"Advanced Tracing (OTEL)","desc":"Explore manual context propagation, custom decorators, and sampling techniques for real-world async, multi-service, and high-volume tracing scenarios.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/advanced-tracing-examples","cat":"Docs"},{"title":"FI Semantic Conventions","desc":"Use standardized attribute keys for spans to ensure consistent, queryable trace data across LLM models, frameworks, and vendors.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/semantic-conventions","cat":"Docs"},{"title":"In-line Evaluations","desc":"Run evaluations directly inside a traced span so results are automatically attached to that span in the Future AGI dashboard.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/in-line-evals","cat":"Docs"},{"title":"Adding Annotations to your Spans","desc":"Label spans with custom tags, human feedback, and notes using the Future AGI bulk-annotation API for systematic trace enrichment.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/annotating-using-api","cat":"Docs"},{"title":"Langfuse Integration","desc":"Integrate Future AGI evaluations with Langfuse to attach evaluation scores and results directly to your Langfuse traces.","href":"https://docs.futureagi.com/docs/observe/features/manual-tracing/langfuse-integration","cat":"Docs"},{"title":"Overview","desc":"Auto-instrumentation integrations for LLM frameworks in Python, JavaScript, and Java. Install a traceAI package to start capturing spans automatically.","href":"https://docs.futureagi.com/docs/tracing/auto","cat":"Docs"},{"title":"OpenAI","desc":"Set up auto-instrumentation for OpenAI with Future AGI tracing. Install traceAI-openai to capture chat completion, embedding, and tool call spans.","href":"https://docs.futureagi.com/docs/tracing/auto/openai","cat":"Docs"},{"title":"Anthropic","desc":"Set up auto-instrumentation for Anthropic Claude with Future AGI tracing. Install traceAI-anthropic to capture LLM spans, inputs, and outputs.","href":"https://docs.futureagi.com/docs/tracing/auto/anthropic","cat":"Docs"},{"title":"AWS Bedrock","desc":"Set up auto-instrumentation for AWS Bedrock with Future AGI tracing. Install traceAI-bedrock to capture model invocation spans and metadata.","href":"https://docs.futureagi.com/docs/tracing/auto/bedrock","cat":"Docs"},{"title":"Vertex AI","desc":"Set up auto-instrumentation for Vertex AI with Future AGI tracing. Install traceAI-vertexai to capture Gemini model invocation and response spans.","href":"https://docs.futureagi.com/docs/tracing/auto/vertexai","cat":"Docs"},{"title":"Google GenAI","desc":"Set up auto-instrumentation for Google GenAI with Future AGI tracing. Install traceAI-google-genai to capture Gemini model interaction spans.","href":"https://docs.futureagi.com/docs/tracing/auto/google_genai","cat":"Docs"},{"title":"Google ADK","desc":"Set up auto-instrumentation for Google ADK with Future AGI tracing. Install traceai-google-adk to capture agent and tool execution spans.","href":"https://docs.futureagi.com/docs/tracing/auto/google_adk","cat":"Docs"},{"title":"Groq","desc":"Set up auto-instrumentation for Groq with Future AGI tracing. Install traceAI-groq to capture high-speed inference spans and performance data.","href":"https://docs.futureagi.com/docs/tracing/auto/groq","cat":"Docs"},{"title":"MistralAI","desc":"Set up auto-instrumentation for Mistral AI with Future AGI tracing. Install traceAI-mistralai to capture model inference spans and metadata.","href":"https://docs.futureagi.com/docs/tracing/auto/mistralai","cat":"Docs"},{"title":"Together AI","desc":"Set up auto-instrumentation for Together AI with Future AGI tracing. Use traceAI-openai to capture inference spans from Together AI models.","href":"https://docs.futureagi.com/docs/tracing/auto/togetherai","cat":"Docs"},{"title":"Ollama","desc":"Set up auto-instrumentation for Ollama with Future AGI tracing. Use traceAI-openai to capture spans from Ollama's OpenAI-compatible local LLM API.","href":"https://docs.futureagi.com/docs/tracing/auto/ollama","cat":"Docs"},{"title":"Portkey","desc":"Set up auto-instrumentation for Portkey with Future AGI tracing. Install traceAI-portkey to capture routed LLM call spans and gateway metrics.","href":"https://docs.futureagi.com/docs/tracing/auto/portkey","cat":"Docs"},{"title":"LangChain","desc":"Set up auto-instrumentation for LangChain with Future AGI tracing. Install traceAI-langchain to capture chain, tool, and LLM call spans.","href":"https://docs.futureagi.com/docs/tracing/auto/langchain","cat":"Docs"},{"title":"LangGraph","desc":"Set up auto-instrumentation for LangGraph with Future AGI tracing. Capture agent graph execution and state transition spans via LangChain instrumentor.","href":"https://docs.futureagi.com/docs/tracing/auto/langgraph","cat":"Docs"},{"title":"LlamaIndex","desc":"Set up auto-instrumentation for LlamaIndex with Future AGI tracing. Install traceAI-llamaindex to capture query, retrieval, and response spans.","href":"https://docs.futureagi.com/docs/tracing/auto/llamaindex","cat":"Docs"},{"title":"LlamaIndex Workflows","desc":"Set up auto-instrumentation for LlamaIndex Workflows with Future AGI tracing. Trace workflow agent execution via the LlamaIndex instrumentor.","href":"https://docs.futureagi.com/docs/tracing/auto/llamaindex-workflows","cat":"Docs"},{"title":"LiteLLM","desc":"Set up auto-instrumentation for LiteLLM with Future AGI tracing. Install traceAI-litellm to capture spans across multiple LLM provider calls.","href":"https://docs.futureagi.com/docs/tracing/auto/litellm","cat":"Docs"},{"title":"CrewAI","desc":"Set up auto-instrumentation for CrewAI with Future AGI tracing. Install traceAI-crewai to capture crew task execution and agent interaction spans.","href":"https://docs.futureagi.com/docs/tracing/auto/crewai","cat":"Docs"},{"title":"AutoGen","desc":"Set up auto-instrumentation for Autogen with Future AGI tracing. Install traceAI-autogen to capture multi-agent conversation spans automatically.","href":"https://docs.futureagi.com/docs/tracing/auto/autogen","cat":"Docs"},{"title":"Haystack","desc":"Set up auto-instrumentation for Haystack with Future AGI tracing. Install traceAI-haystack to capture document pipeline and retrieval spans.","href":"https://docs.futureagi.com/docs/tracing/auto/haystack","cat":"Docs"},{"title":"DSPy","desc":"Set up auto-instrumentation for DSPy with Future AGI tracing. Install traceAI-DSPy to capture program compilation and prediction spans automatically.","href":"https://docs.futureagi.com/docs/tracing/auto/dspy","cat":"Docs"},{"title":"OpenAI Agents","desc":"Set up auto-instrumentation for OpenAI Agents SDK with Future AGI tracing. Install traceAI-openai-agents to capture agent workflow spans.","href":"https://docs.futureagi.com/docs/tracing/auto/openai_agents","cat":"Docs"},{"title":"Smol Agents","desc":"Set up auto-instrumentation for Smol Agents with Future AGI tracing. Install traceAI-smolagents to capture lightweight agent execution spans.","href":"https://docs.futureagi.com/docs/tracing/auto/smol_agents","cat":"Docs"},{"title":"Instructor","desc":"Set up auto-instrumentation for Instructor with Future AGI tracing. Install traceAI-instructor to capture structured output extraction spans.","href":"https://docs.futureagi.com/docs/tracing/auto/instructor","cat":"Docs"},{"title":"PromptFlow","desc":"Set up auto-instrumentation for Prompt Flow with Future AGI tracing. Use traceAI-openai to capture prompt flow execution and LLM call spans.","href":"https://docs.futureagi.com/docs/tracing/auto/promptflow","cat":"Docs"},{"title":"Guardrails","desc":"Set up auto-instrumentation for Guardrails AI with Future AGI tracing. Install traceAI-guardrails to trace validation and LLM interaction spans.","href":"https://docs.futureagi.com/docs/tracing/auto/guardrails","cat":"Docs"},{"title":"MCP","desc":"Set up auto-instrumentation for MCP with Future AGI tracing. Install traceAI-mcp to capture Model Context Protocol server and tool call spans.","href":"https://docs.futureagi.com/docs/tracing/auto/mcp","cat":"Docs"},{"title":"Mastra","desc":"Set up auto-instrumentation for Mastra with Future AGI tracing. Configure @traceai/mastra to export TypeScript agent spans to Future AGI.","href":"https://docs.futureagi.com/docs/tracing/auto/mastra","cat":"Docs"},{"title":"Vercel AI SDK","desc":"Set up auto-instrumentation for Vercel AI SDK with Future AGI tracing. Install @traceai/vercel to capture AI function call spans in Next.js apps.","href":"https://docs.futureagi.com/docs/tracing/auto/vercel","cat":"Docs"},{"title":"LiveKit","desc":"Integrate LiveKit with Future AGI for voice agent observability. Trace real-time voice interactions and monitor agent performance with traceAI-livekit.","href":"https://docs.futureagi.com/docs/tracing/auto/livekit","cat":"Docs"},{"title":"Pipecat","desc":"Set up auto-instrumentation for Pipecat voice apps with Future AGI tracing. Install traceAI-pipecat to capture voice pipeline and processing spans.","href":"https://docs.futureagi.com/docs/tracing/auto/pipecat","cat":"Docs"},{"title":"Overview","desc":"Set up TraceAI for Java applications. Initialize the tracer, configure credentials, and instrument your LLM clients, vector databases, and frameworks.","href":"https://docs.futureagi.com/docs/tracing/auto/java","cat":"Docs"},{"title":"Spring Boot","desc":"Add tracing to Spring Boot apps with Spring AI. Configure application.yml, wrap your ChatModel and EmbeddingModel, and traces are collected automatically.","href":"https://docs.futureagi.com/docs/tracing/auto/spring-boot","cat":"Docs"},{"title":"OpenAI","desc":"Trace OpenAI chat completions, embeddings, and streaming responses in Java with TracedOpenAIClient. Part of the Future AGI Java observability SDK.","href":"https://docs.futureagi.com/docs/tracing/auto/java/openai","cat":"Docs"},{"title":"Anthropic","desc":"Trace Anthropic Messages API calls in Java with TracedAnthropicClient. Uses reflection for cross-version compatibility with the Future AGI Java SDK.","href":"https://docs.futureagi.com/docs/tracing/auto/java/anthropic","cat":"Docs"},{"title":"AWS Bedrock","desc":"Trace AWS Bedrock model invocations in Java with TracedBedrockRuntimeClient. Supports both InvokeModel (raw JSON) and Converse (typed API).","href":"https://docs.futureagi.com/docs/tracing/auto/java/bedrock","cat":"Docs"},{"title":"Cohere","desc":"Trace Cohere chat, embedding, and reranking operations in Java with TracedCohereClient. Part of the Future AGI Java SDK for LLM observability.","href":"https://docs.futureagi.com/docs/tracing/auto/java/cohere","cat":"Docs"},{"title":"Pinecone","desc":"Trace Pinecone vector operations in Java with TracedPineconeIndex. Query, upsert, delete, and fetch with full span instrumentation.","href":"https://docs.futureagi.com/docs/tracing/auto/java/pinecone","cat":"Docs"},{"title":"LLM Providers","desc":"Trace Google GenAI, Vertex AI, Azure OpenAI, Ollama, and Watsonx in Java with Future AGI. All providers use the same TracedClient wrapper pattern.","href":"https://docs.futureagi.com/docs/tracing/auto/java/llm-providers","cat":"Docs"},{"title":"Vector Databases","desc":"Trace vector database operations in Java. Qdrant, Milvus, ChromaDB, Weaviate, MongoDB, Redis, pgvector, Azure AI Search, and Elasticsearch.","href":"https://docs.futureagi.com/docs/tracing/auto/java/vector-databases","cat":"Docs"},{"title":"Frameworks","desc":"Trace LangChain4j and Semantic Kernel operations in Java. Framework-level wrappers that instrument chains, agents, and prompt invocations.","href":"https://docs.futureagi.com/docs/tracing/auto/java/frameworks","cat":"Docs"},{"title":"n8n","desc":"Dynamically retrieve prompts from your Future AGI account in n8n, select specific versions, and compile prompts with variables in the n8n interface.","href":"https://docs.futureagi.com/docs/integrations/traceai/n8n","cat":"Docs"},{"title":"Overview","desc":"Iteratively improve prompts using evaluation-driven feedback and optimization algorithms for higher-quality, more consistent AI responses.","href":"https://docs.futureagi.com/docs/optimization","cat":"Docs"},{"title":"Understanding Optimization","desc":"Explains how prompt optimization works in Future AGI: the feedback loop, key components, algorithms, and how to choose the right one.","href":"https://docs.futureagi.com/docs/optimization/concepts/concept","cat":"Docs"},{"title":"Bayesian Search","desc":"Use Bayesian optimization for few-shot prompt tuning: learns from trials to pick better example sets and configurations.","href":"https://docs.futureagi.com/docs/optimization/optimizers/bayesian-search","cat":"Docs"},{"title":"Meta-Prompt","desc":"The Meta-Prompt optimizer uses a teacher LLM for deep reasoning-based prompt refinement through systematic failure analysis and rewriting.","href":"https://docs.futureagi.com/docs/optimization/optimizers/meta-prompt","cat":"Docs"},{"title":"ProTeGi","desc":"ProTeGi improves prompts by identifying failures, generating critiques, and applying targeted fixes using Textual Gradients for systematic optimization.","href":"https://docs.futureagi.com/docs/optimization/optimizers/protegi","cat":"Docs"},{"title":"PromptWizard","desc":"Learn about PromptWizard, a multi-stage feedback-driven optimizer that improves prompts through a cycle of mutation, critique, and refinement.","href":"https://docs.futureagi.com/docs/optimization/optimizers/promptwizard","cat":"Docs"},{"title":"GEPA","desc":"GEPA (Genetic Pareto) is an evolutionary algorithm that evolves prompts over generations using reflection and mutation for complex optimization.","href":"https://docs.futureagi.com/docs/optimization/optimizers/gepa","cat":"Docs"},{"title":"Random Search","desc":"Random Search is a simple, gradient-free method for establishing a baseline in prompt optimization by exploring random prompt variations.","href":"https://docs.futureagi.com/docs/optimization/optimizers/random-search","cat":"Docs"},{"title":"Using Python SDK","desc":"Run prompt optimization from code using the agent-opt Python library. Configure datasets, optimizers, and evaluation templates programmatically.","href":"https://docs.futureagi.com/docs/optimization/features/using-python-sdk","cat":"Docs"},{"title":"Using Platform","desc":"Run prompt optimization from the Future AGI UI: pick a dataset and column, configure prompt and evals, run optimization, and apply the best prompt.","href":"https://docs.futureagi.com/docs/optimization/features/using-platform","cat":"Docs"},{"title":"Overview","desc":"Create, manage, version, and optimize AI prompts in the Prompt Workbench for reliable and consistent language model outputs.","href":"https://docs.futureagi.com/docs/prompt","cat":"Docs"},{"title":"Prompt Engineering","desc":"What prompt engineering is, how to think about crafting effective prompts, and how the Prompt Workbench supports the iteration process.","href":"https://docs.futureagi.com/docs/prompt/concepts/prompt-engineering","cat":"Docs"},{"title":"Understanding Prompts","desc":"Explains what a prompt is, how it is structured, how variables work, and how prompts connect to models in the Prompt Workbench.","href":"https://docs.futureagi.com/docs/prompt/concepts/understanding-prompts","cat":"Docs"},{"title":"Versions and Labels","desc":"Explains how prompt versioning and deployment labels work in the Prompt Workbench and how to manage multiple versions in Future AGI.","href":"https://docs.futureagi.com/docs/prompt/concepts/versions-and-labels","cat":"Docs"},{"title":"Create Prompt from Scratch","desc":"Build a new prompt manually in the Prompt Workbench with full control over structure, model, parameters, and variables.","href":"https://docs.futureagi.com/docs/prompt/features/create-from-scratch","cat":"Docs"},{"title":"Create from Existing Template","desc":"Start from a pre-built prompt template in the Future AGI Prompt Workbench and customize it for your specific use case and model.","href":"https://docs.futureagi.com/docs/prompt/features/create-from-template","cat":"Docs"},{"title":"Create with AI","desc":"Generate a new prompt from a plain-language description using the Generate with AI feature in the Future AGI Prompt Workbench.","href":"https://docs.futureagi.com/docs/prompt/features/create-with-ai","cat":"Docs"},{"title":"Prompt Workbench Using SDK","desc":"Create, version, and run prompt templates programmatically using the Future AGI SDK for TypeScript/JavaScript or Python applications.","href":"https://docs.futureagi.com/docs/prompt/features/sdk","cat":"Docs"},{"title":"Linked Traces","desc":"Associate prompts with production traces to monitor latency, token usage, and cost per prompt version in the Prompt Workbench.","href":"https://docs.futureagi.com/docs/prompt/features/linked-traces","cat":"Docs"},{"title":"Manage Folders","desc":"Organize prompt templates into folders in the Future AGI Prompt Workbench to keep your workspace navigable as your library grows.","href":"https://docs.futureagi.com/docs/prompt/features/folders","cat":"Docs"},{"title":"Overview","desc":"Future AGI Protect brings real-time safety and policy enforcement directly into your GenAI application flow to prevent harmful outputs.","href":"https://docs.futureagi.com/docs/protect","cat":"Docs"},{"title":"Use Cases","desc":"Future AGI Protect safeguards AI applications with real-time guardrails for security, reliability, and compliance across text, image, and audio modalities.","href":"https://docs.futureagi.com/docs/protect/concepts/concept","cat":"Docs"},{"title":"Run Protect via SDK","desc":"Set up and configure Future AGI Protect to apply real-time safety checks to your AI application's inputs and outputs using the SDK.","href":"https://docs.futureagi.com/docs/protect/features/run-protect","cat":"Docs"},{"title":"Overview","desc":"Test and compare LLM configurations, prompts, and parameters in Future AGI Prototype before deploying changes to production.","href":"https://docs.futureagi.com/docs/prototype","cat":"Docs"},{"title":"Understanding Prototype","desc":"Explains what Prototype is, the problem it solves, and how versions, traces, and evals work together before you ship to production.","href":"https://docs.futureagi.com/docs/prototype/concepts/understanding-prototype","cat":"Docs"},{"title":"Versions and Runs","desc":"What a version is in Prototype, how runs get tagged to a version, and how the dashboard uses versions to compare configurations.","href":"https://docs.futureagi.com/docs/prototype/concepts/versions-and-runs","cat":"Docs"},{"title":"Set Up Prototype","desc":"Configure your environment, register your prototype project, and instrument your app so traces and evals appear in the Prototype dashboard.","href":"https://docs.futureagi.com/docs/prototype/features/set-up-prototype","cat":"Docs"},{"title":"Evals","desc":"Define which evaluations run on your prototype outputs using EvalTags, mapping, and optional custom evals in Future AGI Prototype.","href":"https://docs.futureagi.com/docs/prototype/features/evals","cat":"Docs"},{"title":"Choose Winner","desc":"Rank prototype versions by evaluation scores, cost, and latency, then select and promote the best-performing version to production.","href":"https://docs.futureagi.com/docs/prototype/features/choose-winner","cat":"Docs"},{"title":"Admin & Settings","desc":"Access and manage your Future AGI API keys and secret keys from the developer dashboard for SDK and API authentication.","href":"https://docs.futureagi.com/docs/admin-settings","cat":"Docs"},{"title":"API Keys","desc":"Create, copy, and rotate FI_API_KEY and FI_SECRET_KEY credentials for authenticating with Future AGI Python and TypeScript SDKs.","href":"https://docs.futureagi.com/docs/admin-settings/api-keys","cat":"Docs"},{"title":"Profile & Security","desc":"Update your name, reset your password, enable TOTP two-factor authentication, register passkeys, and manage recovery codes in Future AGI.","href":"https://docs.futureagi.com/docs/admin-settings/profile-security","cat":"Docs"},{"title":"Organization Settings","desc":"Configure your organization's display name and enforce mandatory two-factor authentication for all members in Future AGI.","href":"https://docs.futureagi.com/docs/admin-settings/organization-settings","cat":"Docs"},{"title":"User Management","desc":"Invite team members, assign Owner, Admin, Member, or Viewer roles, manage workspace access, and deactivate users in Future AGI.","href":"https://docs.futureagi.com/docs/admin-settings/user-management","cat":"Docs"},{"title":"Workspace Management","desc":"Create workspaces to isolate environments and teams, manage workspace members, AI providers, integrations, and usage per workspace.","href":"https://docs.futureagi.com/docs/admin-settings/workspace-management","cat":"Docs"},{"title":"AI Providers","desc":"Connect OpenAI, Anthropic, AWS Bedrock, Azure OpenAI, and custom model endpoints to Future AGI for evaluations and optimization.","href":"https://docs.futureagi.com/docs/admin-settings/ai-providers","cat":"Docs"},{"title":"Integrations","desc":"Connect Future AGI to Datadog, PostHog, PagerDuty, Langfuse, S3, and other platforms for observability, alerting, and log archival.","href":"https://docs.futureagi.com/docs/admin-settings/integrations","cat":"Docs"},{"title":"Usage Summary","desc":"Monitor API call counts, input/output token usage, and evaluation runs broken down by month and workspace in Future AGI.","href":"https://docs.futureagi.com/docs/admin-settings/usage-summary","cat":"Docs"},{"title":"Billing & Pricing","desc":"Manage your Future AGI subscription plan, add wallet funds, configure auto-reload, update billing info, and view invoice history.","href":"https://docs.futureagi.com/docs/admin-settings/billing-pricing","cat":"Docs"},{"title":"Roles & Permissions","desc":"Future AGI provides role-based access control (RBAC) with organization roles (Owner, Admin, Member, Viewer) and workspace roles for team access management.","href":"https://docs.futureagi.com/docs/roles-and-permissions","cat":"Docs"},{"title":"Installation","desc":"Install the Future AGI Python SDK and configure your API key and project settings to start evaluating and monitoring AI models.","href":"https://docs.futureagi.com/docs/installation","cat":"Docs"},{"title":"FAQ","desc":"Find answers to frequently asked questions about Future AGI products, pricing, features, SDK setup, and platform usage.","href":"https://docs.futureagi.com/docs/faq","cat":"Docs"},{"title":"Overview","desc":"Test AI agents and prompts through controlled simulations before deploying to production. Run voice and chat simulations, score results, and iterate.","href":"https://docs.futureagi.com/docs/simulation","cat":"Docs"},{"title":"Agent Definition","desc":"An agent definition configures how your AI agent behaves during voice or chat conversations in Future AGI simulation tests.","href":"https://docs.futureagi.com/docs/simulation/concepts/agent-definition","cat":"Docs"},{"title":"Scenarios","desc":"Scenarios defines the test cases, customer profiles, and conversation flows that your AI agent will encounter during simulations.","href":"https://docs.futureagi.com/docs/simulation/concepts/scenarios","cat":"Docs"},{"title":"Personas","desc":"Personas in Future AGI represent the customers or users your AI agent interacts with in simulation. Define them to create realistic test conversations.","href":"https://docs.futureagi.com/docs/simulation/concepts/personas","cat":"Docs"},{"title":"Global Nodes","desc":"A global node is a conversation step the agent can enter at any point. Use it for off-topic questions, interrupts, and 'talk to a human' requests.","href":"https://docs.futureagi.com/docs/simulation/concepts/global-nodes","cat":"Docs"},{"title":"Run Voice Simulation","desc":"Create and run voice simulation tests from the Future AGI platform to evaluate your agent against predefined scenarios and personas.","href":"https://docs.futureagi.com/docs/simulation/features/run-simulation","cat":"Docs"},{"title":"Chat Simulation Using SDK","desc":"Run Future AGI chat simulations from Python by providing an agent callback and executing Run Tests. Automate and scale simulation testing with the SDK.","href":"https://docs.futureagi.com/docs/simulation/features/simulation-using-sdk","cat":"Docs"},{"title":"Chat Replay","desc":"Replay real production sessions in a dev environment using chat simulation to debug, iterate, and improve your agent. Works with Observe data.","href":"https://docs.futureagi.com/docs/simulation/features/observe-to-simulate","cat":"Docs"},{"title":"Voice Replay","desc":"Replay real production voice calls from Future AGI Observe in simulation to debug, iterate, and improve your voice agent based on real usage.","href":"https://docs.futureagi.com/docs/simulation/features/voice-replay","cat":"Docs"},{"title":"Prompt Simulation","desc":"Test your prompts in realistic multi-turn conversations directly from the Prompt Workbench, with no agent deployment or SDK required.","href":"https://docs.futureagi.com/docs/simulation/features/prompt-simulation","cat":"Docs"},{"title":"Evaluate Tool Calling","desc":"Evaluate the tool-calling capabilities of your AI agent during Future AGI simulation runs. Test whether agents invoke the correct tools and parameters.","href":"https://docs.futureagi.com/docs/simulation/features/evaluate-tool-calling","cat":"Docs"},{"title":"View Results","desc":"Read simulation results in Future AGI: view conversation transcripts, evaluation scores, performance analytics, and call logs for each agent run.","href":"https://docs.futureagi.com/docs/simulation/features/view-results","cat":"Docs"},{"title":"Fix My Agent","desc":"Diagnose and fix agent performance issues using in-depth analytics from Future AGI simulation results. Get targeted recommendations for each failure type.","href":"https://docs.futureagi.com/docs/simulation/features/fix-my-agent","cat":"Docs"},{"title":"Overview","desc":"Connect Future AGI with your existing AI frameworks, LLM providers, observability tools, and data export destinations for full-stack AI monitoring.","href":"https://docs.futureagi.com/docs/integrations","cat":"Docs"},{"title":"OpenAI","desc":"Integrate OpenAI with Future AGI for auto-instrumented tracing. Capture chat completions, embeddings, and tool calls with traceAI-openai.","href":"https://docs.futureagi.com/docs/integrations/traceai/openai","cat":"Docs"},{"title":"Anthropic","desc":"Integrate Anthropic Claude with Future AGI for auto-instrumented tracing. Install traceAI-anthropic and capture LLM calls with full observability.","href":"https://docs.futureagi.com/docs/integrations/traceai/anthropic","cat":"Docs"},{"title":"AWS Bedrock","desc":"Integrate AWS Bedrock with Future AGI for auto-instrumented tracing. Capture model invocations and monitor performance with traceAI-bedrock.","href":"https://docs.futureagi.com/docs/integrations/traceai/bedrock","cat":"Docs"},{"title":"Vertex AI","desc":"Integrate Vertex AI (Gemini) with Future AGI observability. Trace model calls and monitor performance using traceAI-vertexai instrumentation.","href":"https://docs.futureagi.com/docs/integrations/traceai/vertexai","cat":"Docs"},{"title":"Google GenAI","desc":"Integrate Google GenAI with Future AGI observability. Set up traceAI-google-genai to capture model calls and monitor performance automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/google_genai","cat":"Docs"},{"title":"Google ADK","desc":"Integrate Google ADK with Future AGI for auto-instrumented tracing. Monitor Google AI agent calls and tool usage with traceAI-google-adk.","href":"https://docs.futureagi.com/docs/integrations/traceai/google_adk","cat":"Docs"},{"title":"Groq","desc":"Integrate Groq with Future AGI observability. Set up traceAI-groq to automatically trace high-speed inference calls and monitor LLM performance.","href":"https://docs.futureagi.com/docs/integrations/traceai/groq","cat":"Docs"},{"title":"MistralAI","desc":"Integrate Mistral AI with Future AGI observability. Set up traceAI-mistralai to capture model calls and monitor inference performance automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/mistralai","cat":"Docs"},{"title":"Together AI","desc":"Integrate Together AI with Future AGI observability. Trace inference calls to Together AI models using the traceAI-openai compatible package.","href":"https://docs.futureagi.com/docs/integrations/traceai/togetherai","cat":"Docs"},{"title":"Ollama","desc":"Integrate Ollama with Future AGI observability. Trace locally-hosted LLM calls using the traceAI-openai package with Ollama's OpenAI-compatible API.","href":"https://docs.futureagi.com/docs/integrations/traceai/ollama","cat":"Docs"},{"title":"Portkey","desc":"Integrate Portkey AI gateway with Future AGI observability. Trace routed LLM calls and monitor performance with traceAI-portkey instrumentation.","href":"https://docs.futureagi.com/docs/integrations/traceai/portkey","cat":"Docs"},{"title":"LangChain","desc":"Integrate LangChain with Future AGI for auto-instrumented tracing. Capture chain executions, tool calls, and LLM interactions with traceAI-langchain.","href":"https://docs.futureagi.com/docs/integrations/traceai/langchain","cat":"Docs"},{"title":"LangGraph","desc":"Integrate LangGraph with Future AGI observability. Trace agent graph execution, tool usage, and state transitions using the LangChain instrumentor.","href":"https://docs.futureagi.com/docs/integrations/traceai/langgraph","cat":"Docs"},{"title":"LlamaIndex","desc":"Integrate LlamaIndex with Future AGI observability. Set up traceAI-llamaindex to trace queries, retrieval, and response generation automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/llamaindex","cat":"Docs"},{"title":"LlamaIndex Workflows","desc":"Integrate LlamaIndex Workflows with Future AGI. Trace workflow-based agent execution and data processing using the LlamaIndex instrumentor.","href":"https://docs.futureagi.com/docs/integrations/traceai/llamaindex-workflows","cat":"Docs"},{"title":"LiteLLM","desc":"Integrate LiteLLM with Future AGI observability. Set up traceAI-litellm to trace calls across multiple LLM providers through a unified interface.","href":"https://docs.futureagi.com/docs/integrations/traceai/litellm","cat":"Docs"},{"title":"CrewAI","desc":"Integrate CrewAI with Future AGI observability. Set up traceAI-crewai to trace multi-agent crew task execution and tool usage automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/crewai","cat":"Docs"},{"title":"AutoGen","desc":"Integrate Autogen with Future AGI observability. Set up traceAI-autogen for automatic tracing of multi-agent conversations and workflows.","href":"https://docs.futureagi.com/docs/integrations/traceai/autogen","cat":"Docs"},{"title":"Haystack","desc":"Integrate Haystack with Future AGI observability. Set up traceAI-haystack to trace document processing pipelines and LLM calls automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/haystack","cat":"Docs"},{"title":"DSPy","desc":"Integrate DSPy with Future AGI observability. Set up traceAI-DSPy to automatically trace DSPy program compilation and inference pipelines.","href":"https://docs.futureagi.com/docs/integrations/traceai/dspy","cat":"Docs"},{"title":"OpenAI Agents","desc":"Integrate OpenAI Agents SDK with Future AGI. Trace agent tool calls, handoffs, and reasoning steps automatically with traceAI-openai-agents.","href":"https://docs.futureagi.com/docs/integrations/traceai/openai_agents","cat":"Docs"},{"title":"Smol Agents","desc":"Integrate Smol Agents with Future AGI observability. Set up traceAI-smolagents to trace lightweight agent tool calls and reasoning automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/smol_agents","cat":"Docs"},{"title":"Instructor","desc":"Integrate Instructor with Future AGI observability. Trace structured LLM output extraction and validation automatically using traceAI-instructor.","href":"https://docs.futureagi.com/docs/integrations/traceai/instructor","cat":"Docs"},{"title":"PromptFlow","desc":"Integrate Prompt Flow with Future AGI observability. Trace prompt flow executions and LLM calls automatically using the traceAI-openai package.","href":"https://docs.futureagi.com/docs/integrations/traceai/promptflow","cat":"Docs"},{"title":"Guardrails","desc":"Integrate Guardrails AI with Future AGI observability. Trace guardrail validations and LLM interactions automatically using traceAI-guardrails.","href":"https://docs.futureagi.com/docs/integrations/traceai/guardrails","cat":"Docs"},{"title":"MCP","desc":"Integrate Model Context Protocol (MCP) with Future AGI. Trace MCP server interactions and tool calls with traceAI-mcp auto-instrumentation.","href":"https://docs.futureagi.com/docs/integrations/traceai/mcp","cat":"Docs"},{"title":"Mastra","desc":"Integrate Mastra with Future AGI for TypeScript agent observability. Configure trace export using the @traceai/mastra package for LLM monitoring.","href":"https://docs.futureagi.com/docs/integrations/traceai/mastra","cat":"Docs"},{"title":"Vercel AI SDK","desc":"Integrate Vercel AI SDK with Future AGI. Set up @traceai/vercel for automatic tracing of AI-powered Next.js and Vercel applications.","href":"https://docs.futureagi.com/docs/integrations/traceai/vercel","cat":"Docs"},{"title":"LiveKit","desc":"Integrate LiveKit with Future AGI observability using traceai-livekit. Trace voice agent sessions, audio pipelines, and tool calls automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/livekit","cat":"Docs"},{"title":"Pipecat","desc":"Integrate Pipecat with Future AGI for voice application observability. Trace and monitor voice pipelines with OpenTelemetry-based traceAI-pipecat.","href":"https://docs.futureagi.com/docs/integrations/traceai/pipecat","cat":"Docs"},{"title":"Overview","desc":"Set up TraceAI for Java applications. Initialize the tracer, configure credentials, and instrument your LLM clients, vector databases, and frameworks.","href":"https://docs.futureagi.com/docs/integrations/traceai/java","cat":"Docs"},{"title":"Spring Boot","desc":"Add tracing to Spring Boot apps with Spring AI. Configure application.yml, wrap your ChatModel and EmbeddingModel, and traces are collected automatically.","href":"https://docs.futureagi.com/docs/integrations/traceai/spring-boot","cat":"Docs"},{"title":"OpenAI","desc":"Trace OpenAI chat completions, embeddings, and streaming responses in Java with TracedOpenAIClient for automatic LLM observability in Future AGI.","href":"https://docs.futureagi.com/docs/integrations/traceai/java/openai","cat":"Docs"},{"title":"Anthropic","desc":"Trace Anthropic Messages API calls in Java with TracedAnthropicClient. Uses reflection for cross-version compatibility and automatic span capture.","href":"https://docs.futureagi.com/docs/integrations/traceai/java/anthropic","cat":"Docs"},{"title":"AWS Bedrock","desc":"Trace AWS Bedrock model invocations in Java with TracedBedrockRuntimeClient. Supports both InvokeModel (raw JSON) and Converse (typed API).","href":"https://docs.futureagi.com/docs/integrations/traceai/java/bedrock","cat":"Docs"},{"title":"Cohere","desc":"Trace Cohere chat, embedding, and reranking operations in Java with TracedCohereClient for automatic LLM observability in Future AGI.","href":"https://docs.futureagi.com/docs/integrations/traceai/java/cohere","cat":"Docs"},{"title":"Pinecone","desc":"Trace Pinecone vector operations in Java with TracedPineconeIndex. Query, upsert, delete, and fetch with full span instrumentation.","href":"https://docs.futureagi.com/docs/integrations/traceai/java/pinecone","cat":"Docs"},{"title":"LLM Providers","desc":"Trace Google GenAI, Vertex AI, Azure OpenAI, Ollama, and Watsonx in Java using Future AGI's Traced wrapper pattern for auto-instrumentation.","href":"https://docs.futureagi.com/docs/integrations/traceai/java/llm-providers","cat":"Docs"},{"title":"Vector Databases","desc":"Trace vector database operations in Java. Qdrant, Milvus, ChromaDB, Weaviate, MongoDB, Redis, pgvector, Azure AI Search, and Elasticsearch.","href":"https://docs.futureagi.com/docs/integrations/traceai/java/vector-databases","cat":"Docs"},{"title":"Frameworks","desc":"Trace LangChain4j and Semantic Kernel operations in Java. Framework-level wrappers that instrument chains, agents, and prompt invocations.","href":"https://docs.futureagi.com/docs/integrations/traceai/java/frameworks","cat":"Docs"},{"title":"n8n","desc":"Dynamically retrieve prompts from your Future AGI account in n8n, select specific versions, and compile prompts with variables in the n8n interface.","href":"https://docs.futureagi.com/docs/integrations/traceai/n8n","cat":"Docs"},{"title":"Langfuse","desc":"Pull existing Langfuse traces, spans, and scores into Future AGI automatically to continue evaluation and analysis without re-instrumentation.","href":"https://docs.futureagi.com/docs/integrations/import/langfuse","cat":"Docs"},{"title":"Datadog","desc":"Forward Agent Command Center logs and metrics from Future AGI to Datadog automatically for centralized monitoring and alerting.","href":"https://docs.futureagi.com/docs/integrations/export/datadog","cat":"Docs"},{"title":"PostHog","desc":"Send LLM usage events from Future AGI's Agent Command Center to PostHog for product analytics, user behavior tracking, and event-based insights.","href":"https://docs.futureagi.com/docs/integrations/export/posthog","cat":"Docs"},{"title":"Mixpanel","desc":"Send LLM usage events from Future AGI's Agent Command Center to Mixpanel for product analytics, user behavior tracking, and funnel analysis.","href":"https://docs.futureagi.com/docs/integrations/export/mixpanel","cat":"Docs"},{"title":"PagerDuty","desc":"Route Future AGI alerts to PagerDuty so your on-call team receives notifications when something breaks in your AI agent pipeline.","href":"https://docs.futureagi.com/docs/integrations/export/pagerduty","cat":"Docs"},{"title":"Cloud Storage","desc":"Archive Agent Command Center logs to Amazon S3, Azure Blob Storage, or Google Cloud Storage as compressed JSONL files for long-term retention.","href":"https://docs.futureagi.com/docs/integrations/export/cloud-storage","cat":"Docs"},{"title":"Message Queues","desc":"Stream Agent Command Center logs from Future AGI to Amazon SQS or Google Pub/Sub for real-time processing and downstream event handling.","href":"https://docs.futureagi.com/docs/integrations/export/message-queues","cat":"Docs"},{"title":"Overview","desc":"Practical step-by-step guides for evaluation, optimization, simulation, observability, RAG, and agent testing with Future AGI products.","href":"https://docs.futureagi.com/docs/cookbook","cat":"Docs"},{"title":"Running Your First Eval","desc":"Score LLM outputs for hallucination, toxicity, and custom criteria using local metrics, Future AGI evaluation models, or LLM-as-Judge.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/first-eval","cat":"Docs"},{"title":"Custom Eval Metrics: Write Your Own Evaluation Criteria","desc":"Define LLM quality criteria in plain English, register reusable eval metrics in the FutureAGI dashboard, and run them via SDK with a single evaluate() call.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/custom-eval-metrics","cat":"Docs"},{"title":"Hallucination Detection with Faithfulness & Groundedness","desc":"Catch LLM hallucinations in RAG outputs using faithfulness (local NLI) and groundedness metrics. Combine both in a single evaluate() call.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/hallucination-detection","cat":"Docs"},{"title":"RAG Pipeline Evaluation: Debug Retrieval vs Generation","desc":"Score RAG retrieval and generation quality independently with five metrics in one evaluate() call to pinpoint whether failures are at retrieval or generation.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/rag-evaluation","cat":"Docs"},{"title":"Multimodal Evaluation: Images, Audio, and PDF","desc":"Score image captions, detect AI-generated images, evaluate audio quality and TTS accuracy, and verify OCR output using built-in multimodal eval metrics.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/multimodal-eval","cat":"Docs"},{"title":"Tone, Toxicity, and Bias Detection Evals","desc":"Evaluate LLM outputs for professional tone, harmful content, and demographic bias using is_polite, toxicity, and bias_detection metrics in the FutureAGI SDK.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/tone-toxicity-bias-eval","cat":"Docs"},{"title":"Evaluate Customer Agent Conversations","desc":"Score multi-turn conversations for quality, context retention, query handling, loop detection, and escalation using built-in conversation metrics.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/conversation-eval","cat":"Docs"},{"title":"Dataset SDK: Upload, Evaluate, and Download Results","desc":"Upload a CSV dataset, run batch evaluations for groundedness and toxicity across every row, and download scored results from the SDK.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/batch-eval","cat":"Docs"},{"title":"Async Evaluations for Large-Scale Testing","desc":"Submit fire-and-forget async evaluations, poll job IDs for results, and run 50+ evals in parallel with ThreadPoolExecutor using the FutureAGI Evaluator SDK.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/async-batch-eval","cat":"Docs"},{"title":"Text-to-SQL Evaluation","desc":"Evaluate LLM-generated SQL queries with string comparison, execution-based validation against a live database, and built-in text_to_sql eval metrics.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/text-to-sql-eval","cat":"Docs"},{"title":"Chat Simulation: Run Multi-Persona Conversations via SDK","desc":"Define personas, auto-generate scenarios, run multi-turn conversations via SDK, and diagnose failures with Fix My Agent using Future AGI Chat Simulation.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/chat-simulation-personas","cat":"Docs"},{"title":"Voice Simulation: Define Agents, Personas, and Run Call Tests","desc":"Define voice agents with caller personas, generate test scenarios, run parallel call tests with evaluations, and diagnose failures with Fix My Agent.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/voice-simulation","cat":"Docs"},{"title":"Tool-Calling Agent Simulation with Tracing","desc":"Simulate tool-calling agent conversations, trace every tool invocation as child spans with fi-instrumentation-otel, and inspect results in the Tracing dashboard.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/tool-calling-simulation","cat":"Docs"},{"title":"Simulate from the Prompt Workbench","desc":"Launch multi-turn chat simulations against any saved prompt version directly from the Future AGI Prompts workbench. No SDK or agent definition required.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/prompt-workbench-simulation","cat":"Docs"},{"title":"Create and Manage Datasets from the Dashboard","desc":"Create a dataset, add columns, enter rows manually or via CSV, run evaluations, and export results from the Future AGI dashboard. No code required.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/dataset-management","cat":"Docs"},{"title":"Synthetic Data Generation: Create Test Datasets from a Schema","desc":"Define column schemas with types and categorical distributions, then generate structured test datasets from the Future AGI dashboard. No code required.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/synthetic-data-generation","cat":"Docs"},{"title":"Annotate Datasets with Human-in-the-Loop Workflows","desc":"Create annotation views with categorical, numeric, and text labels, assign annotators, and log annotations programmatically using the Future AGI SDK.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/dataset-annotation","cat":"Docs"},{"title":"Import Datasets from Hugging Face","desc":"Import any public Hugging Face dataset into Future AGI with a single SDK call, run evaluations, and download scored results.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/huggingface-dataset-import","cat":"Docs"},{"title":"Dynamic Dataset Columns: Enrich Rows with AI-Generated Data","desc":"Enrich datasets with AI-generated summaries, sentiment labels, entities, vector-retrieved context, and parsed JSON fields from the Future AGI dashboard.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/dynamic-dataset-columns","cat":"Docs"},{"title":"Prompt Versioning: Create, Label, and Serve Prompt Versions","desc":"Create prompt templates, commit numbered versions, assign labels like production, and serve the right version at runtime via SDK or dashboard.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/prompt-versioning","cat":"Docs"},{"title":"Prototype and Iterate on LLM Applications","desc":"Register a Prototype project with automatic span evaluation, iterate with versioned prompts, and compare versions before deploying to production.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/prototype-llm-app","cat":"Docs"},{"title":"Manual Tracing: Add Custom Spans to Any Application","desc":"Instrument any Python application with custom spans, user context, and metadata and see every call visualized in the Future AGI Tracing dashboard.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/manual-tracing","cat":"Docs"},{"title":"Session-Based Observability for Multi-Turn Conversations","desc":"Group every span from a multi-turn chatbot by session and user ID so conversations appear as a single, filterable unit in the Future AGI Tracing dashboard.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/session-observability","cat":"Docs"},{"title":"Monitoring & Alerts: Track LLM Performance and Set Quality Thresholds","desc":"Instrument a multi-step RAG agent, analyze latency and cost trends in Charts, and configure warning and critical alerts with email or Slack notifications.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/monitoring-alerts","cat":"Docs"},{"title":"Inline Evals in Tracing: Score Every Response as It's Generated","desc":"Attach quality scores directly to production traces and see faithfulness, toxicity, and custom evals alongside every LLM call in Future AGI Tracing.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/inline-evals-tracing","cat":"Docs"},{"title":"Distributed Tracing: Connect Spans Across Services","desc":"Propagate OpenTelemetry trace context across microservices using W3C TraceContext headers. Spans from gateway to LLM backend land in one unified trace.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/distributed-tracing","cat":"Docs"},{"title":"Prompt Optimization: Improve a Prompt Automatically","desc":"Use agent-opt to take a weak baseline prompt, run automated optimization, and extract the best-performing variant. No manual prompt engineering required.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/prompt-optimization","cat":"Docs"},{"title":"Compare Optimization Strategies: ProTeGi, GEPA, and PromptWizard","desc":"Run three optimization algorithms on the same task with different evaluation metrics and compare results to pick the best strategy for your use case.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/compare-optimizers","cat":"Docs"},{"title":"Dataset Optimization: Improve Prompts Directly in Your Dataset","desc":"Run automated prompt optimization from the dashboard Optimization tab on any Run Prompt column. Review trial results and promote the winner.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/dataset-optimization","cat":"Docs"},{"title":"Protect: Add Safety Guardrails to LLM Outputs","desc":"Screen text for prompt injection, PII, toxicity, and bias with a single Protect API call. Stack multiple safety rules and get structured pass/fail results.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/protect-guardrails","cat":"Docs"},{"title":"Knowledge Base: Upload Documents and Query with the SDK","desc":"Upload documents to a Knowledge Base via dashboard or SDK, and use them for domain-grounded evaluations and synthetic data generation.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/knowledge-base","cat":"Docs"},{"title":"Experimentation: Compare Prompts and Models on a Dataset","desc":"Test multiple prompt variants and models on the same dataset, evaluate outputs, and pick the best configuration using weighted metric comparison.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/experimentation-compare-prompts","cat":"Docs"},{"title":"Evaluation-Driven Development: Score Every Prompt Change Before Shipping","desc":"Build a local eval loop that scores prompts against a test suite, compare before-and-after results, and gate promotion on quality thresholds.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/eval-driven-dev","cat":"Docs"},{"title":"CI/CD Eval Pipeline: Automate Quality Gates in GitHub Actions","desc":"Run automated faithfulness and toxicity eval gates on every pull request using Future AGI. Block merges when scores fall below configured thresholds.","href":"https://docs.futureagi.com/docs/cookbook/quickstart/cicd-eval-pipeline","cat":"Docs"},{"title":"Test and Fix Your Chat Agent with Simulated Conversations","desc":"Simulate multi-turn conversations against your chat agent, evaluate quality automatically, diagnose failure patterns, and optimize the prompt.","href":"https://docs.futureagi.com/docs/cookbook/use-cases/end-to-end-agent-testing","cat":"Docs"},{"title":"Monitor LLM Quality in Production and Catch Regressions","desc":"Auto-score every production LLM response, set alerts for quality regressions, and diagnose failure patterns before users notice with FutureAGI observability.","href":"https://docs.futureagi.com/docs/cookbook/use-cases/production-quality-monitoring","cat":"Docs"},{"title":"End-to-End with Falcon AI: Trace → Debug → Evaluate → Dataset → Fix in One Workflow","desc":"Find failing traces, lock them as a regression set, score the baseline, get a paste-ready prompt diff, and verify the scores recover, all from one Falcon AI conversation.","href":"https://docs.futureagi.com/docs/cookbook/falcon-ai/end-to-end","cat":"Docs"},{"title":"Context-Aware Trace Debugging with Falcon AI","desc":"Falcon AI auto-attaches the failing trace you're viewing, so you can debug it conversationally and get a paste-ready prompt fix without copy-pasting trace IDs.","href":"https://docs.futureagi.com/docs/cookbook/falcon-ai/context-aware-debugging","cat":"Docs"},{"title":"Building Golden Datasets from Production Traces with Falcon AI","desc":"Turn production traces into a curated, ground-truthed golden dataset in one Falcon AI conversation.","href":"https://docs.futureagi.com/docs/cookbook/falcon-ai/eval-datasets-from-traces","cat":"Docs"},{"title":"Cut LLM Costs 80% With Semantic Caching","desc":"Turn on exact and semantic response caching at the gateway so paraphrased duplicate prompts return cached answers instead of paying for the same call twice.","href":"https://docs.futureagi.com/docs/cookbook/command-center/semantic-caching","cat":"Docs"},{"title":"Debug LLM Traces From Your IDE Using Natural Language MCP Queries","desc":"Connect Future AGI's MCP server to Cursor, Claude Code, or VS Code, then debug failing traces, run evals, and annotate spans without leaving your editor.","href":"https://docs.futureagi.com/docs/cookbook/mcp/debug-traces-from-ide","cat":"Docs"},{"title":"Building an Eval Correction Loop: Teaching Your Evaluator What 'Good' Means for Your Domain","desc":"Run a built-in eval, mark the rows where it disagrees with your judgment, bake those corrections into a custom eval, and re-run until the eval matches how your team scores quality.","href":"https://docs.futureagi.com/docs/cookbook/evaluation/eval-correction-loop","cat":"Docs"},{"title":"Deploy the Full Open-Source AI Stack Locally With Docker Compose in 5 Minutes","desc":"Clone the Future AGI repo, configure .env, run `docker compose up`, and start sending traces. Five commands to a complete self-hosted stack on your laptop.","href":"https://docs.futureagi.com/docs/cookbook/self-hosting/docker-compose-quickstart","cat":"Docs"},{"title":"Using FutureAGI Evals","desc":"Evaluate AI model outputs using the ai-evaluation package. Choose built-in templates, pass inputs, and receive scored results with the Evaluator class.","href":"https://docs.futureagi.com/docs/cookbook/using-futureagi-evals","cat":"Docs"},{"title":"Using FutureAGI Protect","desc":"Apply guardrail rules to LLM outputs with the Future AGI Protect class. Define tone, toxicity, and custom metric rules to block or flag harmful responses.","href":"https://docs.futureagi.com/docs/cookbook/using-futureagi-protect","cat":"Docs"},{"title":"Using FutureAGI Dataset","desc":"Create and manage AI evaluation datasets using the Future AGI Python SDK. Define schema columns, add rows, and access datasets via the futureagi package.","href":"https://docs.futureagi.com/docs/cookbook/using-futureagi-dataset","cat":"Docs"},{"title":"Using FutureAGI KB","desc":"Initialize a KnowledgeBase client with the Future AGI SDK to create and manage knowledge bases, add files, update contents, and query stored documents.","href":"https://docs.futureagi.com/docs/cookbook/using-futureagi-kb","cat":"Docs"},{"title":"Portkey Integration","desc":"Combine Portkey and Future AGI for end-to-end LLM observability. Benchmark multiple models on response quality, latency, and cost.","href":"https://docs.futureagi.com/docs/cookbook/portkey-integration","cat":"Docs"},{"title":"LangChain/LangGraph","desc":"Add observability and evaluation to LangChain and LangGraph agents using Future AGI tracing. Detect completeness, groundedness, and hallucinations.","href":"https://docs.futureagi.com/docs/cookbook/langchain-langgraph","cat":"Docs"},{"title":"LlamaIndex PDF RAG","desc":"Build a production-ready LlamaIndex PDF RAG chatbot with Future AGI observability, tracing, and real-time evaluation of retrieval quality.","href":"https://docs.futureagi.com/docs/cookbook/llamaindex-pdf-rag","cat":"Docs"},{"title":"CrewAI Research Team","desc":"Build a multi-agent research system using CrewAI with Future AGI observability and in-line evaluations for real-time quality monitoring.","href":"https://docs.futureagi.com/docs/cookbook/crewai-research-team","cat":"Docs"},{"title":"MongoDB","desc":"Build a PDF RAG chatbot using MongoDB Atlas for vector search and Future AGI for tracing, evaluation, and LLM pipeline monitoring.","href":"https://docs.futureagi.com/docs/cookbook/mongodb","cat":"Docs"},{"title":"Meeting Summarization","desc":"Evaluate meeting summarization quality using Future AGI. Score AI-generated summaries from transcripts for accuracy and completeness.","href":"https://docs.futureagi.com/docs/cookbook/meeting-summarization","cat":"Docs"},{"title":"AI SDR Evaluation","desc":"Evaluate AI-generated sales outreach messages using Future AGI. Score SDR openers for relevance, personalization, and value proposition alignment.","href":"https://docs.futureagi.com/docs/cookbook/ai-sdr","cat":"Docs"},{"title":"AI Agents Evaluation","desc":"Evaluate AI agent function-calling and response quality using Future AGI's evaluation SDK with metrics like tool use accuracy and safety.","href":"https://docs.futureagi.com/docs/cookbook/ai-agents","cat":"Docs"},{"title":"Image Evaluation","desc":"Evaluate AI-generated images for description alignment, artistic requirements, and replacement quality using the Future AGI SDK.","href":"https://docs.futureagi.com/docs/cookbook/image-evaluation","cat":"Docs"},{"title":"Implement Observability","desc":"Instrument a LangChain chatbot with Future AGI tracing to monitor LLM performance, track metrics, and improve application stability.","href":"https://docs.futureagi.com/docs/cookbook/observability","cat":"Docs"},{"title":"Text-to-SQL Evaluation","desc":"Build and evaluate a Text-to-SQL agent with Future AGI. Test natural language to SQL conversion accuracy using automated evaluation metrics.","href":"https://docs.futureagi.com/docs/cookbook/text-to-sql","cat":"Docs"},{"title":"RAG with LangChain","desc":"Experiment with LangChain RAG configurations using Future AGI. Build and evaluate a retrieval-augmented generation app with OpenAI embeddings.","href":"https://docs.futureagi.com/docs/cookbook/rag-langchain","cat":"Docs"},{"title":"Evaluate RAG Apps","desc":"Evaluate RAG applications with Future AGI using context adherence, retrieval quality, answer correctness, and other retrieval-augmented generation metrics.","href":"https://docs.futureagi.com/docs/cookbook/evaluate-rag","cat":"Docs"},{"title":"Trustworthy RAG Chatbots","desc":"Evaluate RAG chatbot trustworthiness across retrieval accuracy, prompt injection resilience, privacy compliance, and tone adaptation with Future AGI.","href":"https://docs.futureagi.com/docs/cookbook/trustworthy-rag","cat":"Docs"},{"title":"Decrease RAG Hallucination","desc":"Reduce hallucinations in RAG pipelines by benchmarking chunking, retrieval, and chain strategies with Future AGI's evaluation suite.","href":"https://docs.futureagi.com/docs/cookbook/decrease-hallucination","cat":"Docs"},{"title":"End-to-End Prompt Optimization","desc":"Optimize prompts end-to-end with Future AGI. Learn evaluation-driven prompt refinement using automated scoring and version tracking.","href":"https://docs.futureagi.com/docs/cookbook/end-to-end-optimization","cat":"Docs"},{"title":"Basic Prompt Optimization","desc":"Optimize your first prompt using the agent-opt Python library with RandomSearchOptimizer. Generate prompt variations and select the best performer.","href":"https://docs.futureagi.com/docs/cookbook/basic-optimization","cat":"Docs"},{"title":"GEPA Optimization","desc":"A guide to using GEPA, a powerful evolutionary algorithm for state-of-the-art prompt optimization in complex, high-stakes scenarios.","href":"https://docs.futureagi.com/docs/cookbook/gepa-optimization","cat":"Docs"},{"title":"Eval Metrics for Optimization","desc":"Compare three methods for evaluating prompt quality within agent-opt: the Future AGI platform, local LLM-as-a-judge, and local heuristic metrics.","href":"https://docs.futureagi.com/docs/cookbook/eval-metrics-optimization","cat":"Docs"},{"title":"Compare Strategies","desc":"A practical guide to selecting the best optimization strategy (Bayesian Search, Meta-Prompt, GEPA, etc.) based on your specific task and goals.","href":"https://docs.futureagi.com/docs/cookbook/compare-optimization","cat":"Docs"},{"title":"Import Datasets","desc":"Prepare and integrate datasets from in-memory, CSV, JSON, and JSONL sources for prompt optimization with the agent-opt library.","href":"https://docs.futureagi.com/docs/cookbook/import-datasets","cat":"Docs"},{"title":"Chat Simulation with Fix My Agent","desc":"Simulate AI chat agents across multiple scenarios, analyze performance metrics, and use Fix My Agent for AI-powered diagnostics.","href":"https://docs.futureagi.com/docs/cookbook/chat-simulation-fix-agent","cat":"Docs"},{"title":"Simulate SDK Demo","desc":"Use the agent-simulate SDK to build automated tests for conversational voice AI agents and validate agent behavior across call scenarios.","href":"https://docs.futureagi.com/docs/cookbook/simulate-sdk","cat":"Docs"},{"title":"Error Feed with Google ADK","desc":"Build a Google ADK multi-agent system with tracing, then use Future AGI Error Feed to analyze agent errors and surface recommendations on each trace.","href":"https://docs.futureagi.com/docs/cookbook/error-feed/google-adk-multi-agent","cat":"Docs"},{"title":"SDK Overview","desc":"Evaluate LLM outputs, trace AI calls, optimize prompts, and test voice agents. Python, TypeScript, Java, and C# supported.","href":"https://docs.futureagi.com/docs/sdk","cat":"Docs"},{"title":"Overview","desc":"Evaluate LLM outputs with 76+ local metrics, cloud Turing models, or custom LLM-as-Judge criteria. Part of the ai-evaluation Python package.","href":"https://docs.futureagi.com/docs/sdk/evals","cat":"Docs"},{"title":"Running Evaluations","desc":"Run evaluations with the evaluate() function: local heuristics, cloud Turing, or LLM-as-Judge, auto-routed based on your inputs.","href":"https://docs.futureagi.com/docs/sdk/evals/evaluate","cat":"Docs"},{"title":"Distributed Evaluator","desc":"Run evaluations at scale with blocking, async, or distributed execution. Backends include ThreadPool, Celery, Ray, Temporal, and Kubernetes.","href":"https://docs.futureagi.com/docs/sdk/evals/distributed","cat":"Docs"},{"title":"AutoEval","desc":"Auto-generate evaluation pipelines from app descriptions. Pre-built templates for customer support, RAG, code assistants, healthcare, and more.","href":"https://docs.futureagi.com/docs/sdk/evals/autoeval","cat":"Docs"},{"title":"Guardrails","desc":"Screen AI inputs and outputs with model-based safety checks and fast local scanners. 14 guard models, 14 scanners, async and batch support.","href":"https://docs.futureagi.com/docs/sdk/evals/guardrails-module","cat":"Docs"},{"title":"Local & Hybrid","desc":"Run evaluations locally with zero API calls. Auto-route between local and cloud metrics. Use Ollama for offline LLM-based scoring.","href":"https://docs.futureagi.com/docs/sdk/evals/local","cat":"Docs"},{"title":"OpenTelemetry","desc":"Built-in OpenTelemetry for the AI evaluation SDK. Auto-instrument LLM calls, track costs, enrich spans with scores, and export to any backend.","href":"https://docs.futureagi.com/docs/sdk/evals/otel","cat":"Docs"},{"title":"Code Security","desc":"AST-based vulnerability detection for AI-generated code. 15 detectors, 4 evaluation modes, multi-language support, and dual-judge scoring.","href":"https://docs.futureagi.com/docs/sdk/evals/code-security","cat":"Docs"},{"title":"Overview","desc":"Browse all 76+ evaluation metrics by category: string matching, JSON validation, hallucination, RAG quality, agent trajectories, and guardrail checks.","href":"https://docs.futureagi.com/docs/sdk/evals/metrics","cat":"Docs"},{"title":"String & Similarity","desc":"23 local metrics for keyword matching, regex, length checks, BLEU, ROUGE, Levenshtein, and embedding similarity. All run locally via evaluate().","href":"https://docs.futureagi.com/docs/sdk/evals/metrics/string","cat":"Docs"},{"title":"JSON & Structured","desc":"14 metrics for validating JSON correctness, schema compliance, type checking, and structured output quality. Part of the ai-evaluation Python SDK.","href":"https://docs.futureagi.com/docs/sdk/evals/metrics/json","cat":"Docs"},{"title":"Hallucination","desc":"Detect hallucinations, unsupported claims, and contradictions in LLM outputs. 5 context-grounded metrics with optional NLI and LLM augmentation.","href":"https://docs.futureagi.com/docs/sdk/evals/metrics/hallucination","cat":"Docs"},{"title":"RAG","desc":"19 local metrics for evaluating RAG pipelines: retrieval quality, generation faithfulness, advanced reasoning, and composite scores.","href":"https://docs.futureagi.com/docs/sdk/evals/metrics/rag","cat":"Docs"},{"title":"Agents & Functions","desc":"11 metrics for evaluating agent trajectories, tool use, reasoning quality, and function call correctness. All run locally via evaluate().","href":"https://docs.futureagi.com/docs/sdk/evals/metrics/agents","cat":"Docs"},{"title":"Guardrails","desc":"Security-focused scanner metrics that detect prompt injection, PII, secrets, and SQL injection in under 10ms. Part of the ai-evaluation Python SDK.","href":"https://docs.futureagi.com/docs/sdk/evals/metrics/guardrails","cat":"Docs"},{"title":"Cloud Evals","desc":"Run pre-built evaluation templates on Future AGI's Turing cloud models. 100+ templates covering safety, RAG, hallucination, conversation quality, and more.","href":"https://docs.futureagi.com/docs/sdk/evals/cloud-evals","cat":"Docs"},{"title":"LLM-as-Judge","desc":"Define custom grading criteria and run them with any LLM: GPT-4o, Gemini, Claude, Ollama, or any LiteLLM-supported model.","href":"https://docs.futureagi.com/docs/sdk/evals/llm-judge","cat":"Docs"},{"title":"Streaming","desc":"Check LLM output token-by-token as it streams. Detect toxic content, PII, or quality drops mid-generation and stop early.","href":"https://docs.futureagi.com/docs/sdk/evals/streaming","cat":"Docs"},{"title":"Feedback Loops","desc":"Submit corrections to scoring results, calibrate thresholds over time, and store feedback in ChromaDB for continuous improvement.","href":"https://docs.futureagi.com/docs/sdk/evals/feedback","cat":"Docs"},{"title":"Datasets","desc":"Create, populate, and manage datasets for evaluation. Upload CSV/JSON files, import from HuggingFace, add LLM-generated columns, and run evals at scale.","href":"https://docs.futureagi.com/docs/sdk/datasets","cat":"Docs"},{"title":"Tracing","desc":"Set up OpenTelemetry tracing across Python, TypeScript, Java, and C#. Auto-instrument 45+ frameworks or create custom spans with FITracer.","href":"https://docs.futureagi.com/docs/sdk/tracing","cat":"Docs"},{"title":"Protect","desc":"Guard AI inputs and outputs in real-time. Check for content moderation, bias, security threats, and data privacy violations.","href":"https://docs.futureagi.com/docs/sdk/protect","cat":"Docs"},{"title":"Knowledge Base","desc":"Upload documents to build and manage knowledge bases for RAG evaluation and context injection. Create, update, delete, and query files with the Python SDK.","href":"https://docs.futureagi.com/docs/sdk/knowledgebase","cat":"Docs"},{"title":"Annotation Queues","desc":"Reference for the AnnotationQueue class in the Future AGI Python SDK, covering how to create, fetch, and populate annotation queues programmatically.","href":"https://docs.futureagi.com/docs/sdk/annotation-queues","cat":"Docs"},{"title":"Prompt Optimization","desc":"Automatically improve your prompts with 6 SOTA algorithms. Random Search, Bayesian, ProTeGi, Meta-Prompt, PromptWizard, and GEPA.","href":"https://docs.futureagi.com/docs/sdk/optimization","cat":"Docs"},{"title":"Simulation Testing","desc":"Test voice AI agents at scale with simulated customer personas. Run conversations, capture audio, and score agent performance automatically.","href":"https://docs.futureagi.com/docs/sdk/simulate","cat":"Docs"},{"title":"Introduction","desc":"Complete REST API reference for the Future AGI platform. Covers simulations, evaluations, scenarios, personas, datasets, annotations, and more.","href":"https://docs.futureagi.com/docs/api","cat":"Docs"},{"title":"Health Check","desc":"Check whether the Future AGI API server is up and reachable. Returns 200 with a status flag and confirmation message when the server is running.","href":"https://docs.futureagi.com/docs/api/health/healthcheck","cat":"Docs"},{"title":"List Eval Tasks","desc":"Retrieve a paginated list of eval tasks. Filter by project UUID or name. Returns task status, sampling rate, eval count, and last run timestamp.","href":"https://docs.futureagi.com/docs/api/eval-tasks/list-eval-tasks-filtered","cat":"Docs"},{"title":"Create Eval Task","desc":"Create an eval task for a project. Set eval configs, sampling rate, run type (continuous or historical), span filters, and spans limit. Returns the new task UUID.","href":"https://docs.futureagi.com/docs/api/eval-tasks/create-eval-task","cat":"Docs"},{"title":"Get Eval Task","desc":"Retrieve a single eval task by UUID. Returns status, sampling rate, run type, attached eval configs, filters, failed spans, and timestamps.","href":"https://docs.futureagi.com/docs/api/eval-tasks/get-eval-task","cat":"Docs"},{"title":"Update Eval Task","desc":"Partially update an eval task's name, evals, sampling rate, run type, or filters. Use edit_type to choose fresh_run or edit_rerun mode.","href":"https://docs.futureagi.com/docs/api/eval-tasks/update-eval-task","cat":"Docs"},{"title":"Delete Eval Task","desc":"Soft-delete a single eval task by UUID. The task must not be in running state. Returns 204 No Content on success.","href":"https://docs.futureagi.com/docs/api/eval-tasks/delete-eval-task","cat":"Docs"},{"title":"Bulk Delete Eval Tasks","desc":"Soft-delete multiple eval tasks in one request. Accepts an array of eval task UUIDs; tasks must be in a non-running state. Returns success status.","href":"https://docs.futureagi.com/docs/api/eval-tasks/bulk-delete-eval-tasks","cat":"Docs"},{"title":"Pause Eval Task","desc":"Pause a running eval task by UUID. Task must be in running state. Returns a confirmation status and message on success.","href":"https://docs.futureagi.com/docs/api/eval-tasks/pause-eval-task","cat":"Docs"},{"title":"Unpause Eval Task","desc":"Resume a paused eval task by UUID. Task must be in paused state; status resets to pending on resume. Returns confirmation status and message.","href":"https://docs.futureagi.com/docs/api/eval-tasks/unpause-eval-task","cat":"Docs"},{"title":"List Custom Eval Configs","desc":"List all custom eval configs, with optional filtering by project UUID or eval task UUID to narrow results.","href":"https://docs.futureagi.com/docs/api/custom-eval-configs/list-configs-filtered","cat":"Docs"},{"title":"Create Custom Eval Config","desc":"Create a new custom eval config for a project by specifying an eval template, name, field mapping, and optional config settings.","href":"https://docs.futureagi.com/docs/api/custom-eval-configs/create-custom-eval-config","cat":"Docs"},{"title":"Get Custom Eval Config","desc":"Retrieve a specific custom eval config by UUID, returning its template, name, project, field mapping, and current configuration.","href":"https://docs.futureagi.com/docs/api/custom-eval-configs/get-custom-eval-config","cat":"Docs"},{"title":"Update Custom Eval Config","desc":"Partially update an existing custom eval config by UUID. Supports patching name, mapping, config, or template fields individually.","href":"https://docs.futureagi.com/docs/api/custom-eval-configs/update-custom-eval-config","cat":"Docs"},{"title":"Delete Custom Eval Config","desc":"Soft-delete a custom eval config by UUID. Returns 204 No Content on success. The config is deactivated and no longer applied.","href":"https://docs.futureagi.com/docs/api/custom-eval-configs/delete-custom-eval-config","cat":"Docs"},{"title":"Check Config Exists","desc":"Check if a custom eval config with the same name and field mapping already exists in a project before creating a duplicate.","href":"https://docs.futureagi.com/docs/api/custom-eval-configs/check-config-exists","cat":"Docs"},{"title":"Get Eval Template Names","desc":"Search and retrieve evaluation template names available in your organization. Supports text search to filter results by name.","href":"https://docs.futureagi.com/docs/api/dataset-evals/get-eval-template-names","cat":"Docs"},{"title":"Create Custom Eval Template","desc":"Create a reusable custom eval template with criteria, output type (e.g. Pass/Fail), required keys, and model configuration.","href":"https://docs.futureagi.com/docs/api/dataset-evals/create-custom-eval-template","cat":"Docs"},{"title":"List Dataset Evals","desc":"List available evaluations for a dataset. Filter by name, category (built-in or user), type, or tags to find the right eval.","href":"https://docs.futureagi.com/docs/api/dataset-evals/list-dataset-evals","cat":"Docs"},{"title":"Get Eval Structure","desc":"Retrieve the configuration structure of an eval by ID and type (preset, user, or previously configured), including required keys and mapping.","href":"https://docs.futureagi.com/docs/api/dataset-evals/get-eval-structure","cat":"Docs"},{"title":"Add Dataset Eval","desc":"Add an evaluation to a dataset by selecting a template and configuring the input/output key mapping. Returns the created eval record.","href":"https://docs.futureagi.com/docs/api/dataset-evals/add-dataset-eval","cat":"Docs"},{"title":"Start Evals Process","desc":"Trigger one or more evaluations to run across a dataset by submitting user eval IDs. Scores are computed for every row.","href":"https://docs.futureagi.com/docs/api/dataset-evals/start-evals-process","cat":"Docs"},{"title":"Delete Dataset Eval","desc":"Remove an evaluation from a dataset by eval ID. Optionally delete the associated eval column and its data from the dataset.","href":"https://docs.futureagi.com/docs/api/dataset-evals/delete-dataset-eval","cat":"Docs"},{"title":"Edit and Run Eval","desc":"Update an evaluation's config and key mapping, then optionally re-run it across all dataset rows to refresh scores.","href":"https://docs.futureagi.com/docs/api/dataset-evals/edit-and-run-eval","cat":"Docs"},{"title":"List Scenarios","desc":"List paginated scenarios with optional search and filtering by agent definition or agent type. Returns scenario type, status, dataset row count, and creation timestamp.","href":"https://docs.futureagi.com/docs/api/scenarios/listscenarios","cat":"Docs"},{"title":"Get Scenario Details","desc":"Retrieve a scenario by UUID. Returns scenario type, dataset ID, agent type, status, conversation graph, simulator prompts, and dataset row count.","href":"https://docs.futureagi.com/docs/api/scenarios/getscenario","cat":"Docs"},{"title":"Create Scenario","desc":"Create a simulation scenario from a dataset, script, or conversation graph. Supports AI graph generation, persona assignment, and custom columns. Returns 202 Accepted.","href":"https://docs.futureagi.com/docs/api/scenarios/createscenario","cat":"Docs"},{"title":"Edit Scenario","desc":"Update a scenario's name, description, conversation graph, or simulator prompt. Supports persona and situation template variables in the prompt field.","href":"https://docs.futureagi.com/docs/api/scenarios/editscenario","cat":"Docs"},{"title":"Delete Scenario","desc":"Soft-delete a scenario by UUID, marking it as deleted. Returns a confirmation message. Deleted scenarios cannot be recovered through the API.","href":"https://docs.futureagi.com/docs/api/scenarios/deletescenario","cat":"Docs"},{"title":"Add Rows with AI","desc":"Generate and add 10–100 AI-populated rows to a scenario dataset. Provide optional generation guidance; existing rows and columns are used as context. Returns 202 Accepted.","href":"https://docs.futureagi.com/docs/api/scenarios/addscenariorowswithai","cat":"Docs"},{"title":"Add Columns","desc":"Add up to 10 AI-generated columns to a scenario dataset. Specify name, data type, and description per column. Generation is asynchronous; returns 202 Accepted.","href":"https://docs.futureagi.com/docs/api/scenarios/addcolumns","cat":"Docs"},{"title":"Add Empty Rows","desc":"Add a specified number of empty rows to a scenario's dataset by dataset UUID. Accepts num_rows as a positive integer. Returns success status and confirmation message.","href":"https://docs.futureagi.com/docs/api/scenarios/addemptyrowstodataset","cat":"Docs"},{"title":"List Personas","desc":"List system and workspace personas with pagination. Filter by type (prebuilt/custom), simulation type (voice/text), or keyword search. Returns persona details and traits.","href":"https://docs.futureagi.com/docs/api/personas/listpersonas","cat":"Docs"},{"title":"Create Persona","desc":"Create a workspace-level persona for voice or text simulation. Set gender, age group, personality, communication style, and tone. Returns the new persona object.","href":"https://docs.futureagi.com/docs/api/personas/createpersona","cat":"Docs"},{"title":"Update Persona","desc":"Partially update a workspace-level persona's attributes, including name, personality, tone, and communication style. System personas cannot be modified.","href":"https://docs.futureagi.com/docs/api/personas/updatepersona","cat":"Docs"},{"title":"Delete Persona","desc":"Soft-delete a workspace-level persona by UUID. System (prebuilt) personas cannot be deleted. Returns 204 No Content on success.","href":"https://docs.futureagi.com/docs/api/personas/deletepersona","cat":"Docs"},{"title":"Duplicate Persona","desc":"Copy an existing persona (system or workspace) as a new workspace-level persona. Provide a unique name; all other attributes are inherited from the source.","href":"https://docs.futureagi.com/docs/api/personas/duplicatepersona","cat":"Docs"},{"title":"List Agent Definitions","desc":"Return a paginated list of agent definitions. Filter by agent type, search by name or assistant ID, and pin a specific agent to the top of results.","href":"https://docs.futureagi.com/docs/api/agent-definitions/listagentdefinitions","cat":"Docs"},{"title":"Create Agent Definition","desc":"Create a new agent definition and its initial version. Accepts agent type, name, provider, API key, and system prompt. Returns the created agent object.","href":"https://docs.futureagi.com/docs/api/agent-definitions/createagentdefinition","cat":"Docs"},{"title":"Get Agent Definition","desc":"Retrieve a specific agent definition by UUID, including its full version history, provider details, and configuration snapshot.","href":"https://docs.futureagi.com/docs/api/agent-definitions/getagentdefinition","cat":"Docs"},{"title":"Delete Agent Definitions","desc":"Bulk soft-delete one or more agent definitions and all associated versions. Accepts a list of agent UUIDs and returns a count of agents updated.","href":"https://docs.futureagi.com/docs/api/agent-definitions/deleteagentdefinitions","cat":"Docs"},{"title":"Fetch from Provider","desc":"Fetch an assistant's name and system prompt from an external voice provider such as Vapi. Requires assistant ID, API key, and provider name.","href":"https://docs.futureagi.com/docs/api/agent-definitions/fetchassistantfromprovider","cat":"Docs"},{"title":"List Agent Versions","desc":"Retrieve a paginated list of all versions for a given agent definition. Supports limit and page parameters. Returns version metadata and commit history.","href":"https://docs.futureagi.com/docs/api/agent-versions/listagentversions","cat":"Docs"},{"title":"Create Agent Version","desc":"Create a new version of an agent definition. Accepts agent name, language, system prompt, commit message, and provider config. Returns the new version object.","href":"https://docs.futureagi.com/docs/api/agent-versions/createagentversion","cat":"Docs"},{"title":"Get Agent Version","desc":"Retrieve a specific agent version by agent UUID and version UUID, including its full configuration snapshot, prompt, and provider settings.","href":"https://docs.futureagi.com/docs/api/agent-versions/getagentversion","cat":"Docs"},{"title":"Get Version Call Executions","desc":"Retrieve paginated call executions and their evaluation results for a specific agent version. Supports limit and page query parameters.","href":"https://docs.futureagi.com/docs/api/agent-versions/getversioncallexecutions","cat":"Docs"},{"title":"Get Version Eval Summary","desc":"Retrieve aggregated evaluation summary statistics for an agent version, including pass rate, fail count, and error count per eval metric.","href":"https://docs.futureagi.com/docs/api/agent-versions/getversionevalsummary","cat":"Docs"},{"title":"List Test Runs","desc":"List paginated test runs. Filter by name, source type, or prompt template. Returns agent definition, scenarios, tool evaluation flag, and last run timestamp per run.","href":"https://docs.futureagi.com/docs/api/run-tests/listruntests","cat":"Docs"},{"title":"Create Run Test","desc":"Create a new test run with scenarios, agent definition, eval configs, and optional tool evaluation. Returns the run test UUID, name, and associated scenario list.","href":"https://docs.futureagi.com/docs/api/run-tests/createruntest","cat":"Docs"},{"title":"Get Test Run Details","desc":"Retrieve full details of a test run by UUID, including agent definition, prompt template, scenarios, eval configs, tool evaluation settings, and last run timestamp.","href":"https://docs.futureagi.com/docs/api/run-tests/getruntestdetails","cat":"Docs"},{"title":"Delete Test Run","desc":"Soft-delete a test run by UUID. The run must have no currently active executions. Returns a confirmation message on success.","href":"https://docs.futureagi.com/docs/api/run-tests/deleteruntest","cat":"Docs"},{"title":"Execute Run Test","desc":"Trigger a new execution of a test run. Select scenarios by inclusion or exclusion, optionally specify a simulator. Returns execution ID, status, and scenario/call counts.","href":"https://docs.futureagi.com/docs/api/run-tests/executeruntest","cat":"Docs"},{"title":"Update Components","desc":"Update the agent definition, agent version, simulator agent, scenarios, or tool evaluation settings of an existing test run. Returns the updated test run object.","href":"https://docs.futureagi.com/docs/api/run-tests/updatetestcomponents","cat":"Docs"},{"title":"Get Test Executions","desc":"List paginated test executions for a run test. Filter by status or search string. Returns execution status, call counts, success rate, duration, and scenario IDs.","href":"https://docs.futureagi.com/docs/api/run-tests/gettestexecutions","cat":"Docs"},{"title":"Get Test Scenarios","desc":"List paginated scenarios associated with a test run. Supports name search. Returns scenario UUID, display name, and row count per scenario.","href":"https://docs.futureagi.com/docs/api/run-tests/gettestscenarios","cat":"Docs"},{"title":"Get Call Executions","desc":"List paginated call executions for a test run. Filter by status or search string. Returns transcript, eval scores, latency, scenario info, and call summary per call.","href":"https://docs.futureagi.com/docs/api/run-tests/getcallexecutions","cat":"Docs"},{"title":"Get Eval Summary","desc":"Retrieve aggregated evaluation summary for a test run. Optionally scope to a specific execution. Returns per-config average scores, pass counts, and fail counts.","href":"https://docs.futureagi.com/docs/api/run-tests/getevalsummary","cat":"Docs"},{"title":"Compare Eval Summaries","desc":"Compare evaluation summaries side-by-side across multiple test executions. Accepts a JSON-encoded array of execution UUIDs; returns per-execution eval metrics keyed by execution ID.","href":"https://docs.futureagi.com/docs/api/run-tests/compareevalsummaries","cat":"Docs"},{"title":"Add Eval Configs","desc":"Add one or more evaluation configurations to an existing test run. Specify template, name, config mapping, model, and error localizer per config. Returns created eval config objects.","href":"https://docs.futureagi.com/docs/api/run-tests/addevalconfigs","cat":"Docs"},{"title":"Update Eval Config","desc":"Update an evaluation configuration for a test run. Modify config, mapping, model, name, or error localizer. Optionally trigger an immediate rerun after saving.","href":"https://docs.futureagi.com/docs/api/run-tests/updateevalconfig","cat":"Docs"},{"title":"Delete Eval Config","desc":"Delete an evaluation configuration from a test run by run test and eval config UUID. Cannot delete the last remaining config in the run. Returns a confirmation message.","href":"https://docs.futureagi.com/docs/api/run-tests/deleteevalconfig","cat":"Docs"},{"title":"Run New Evals","desc":"Run new evaluation configs on completed test executions. Specify eval config UUIDs and target executions or use selectAll. Returns call execution count being evaluated.","href":"https://docs.futureagi.com/docs/api/run-tests/runnewevalsontestexecution","cat":"Docs"},{"title":"Rerun Test Executions","desc":"Rerun test executions within a run test. Choose eval_only to re-evaluate existing data or call_and_eval to re-execute calls from scratch. Supports selectAll mode.","href":"https://docs.futureagi.com/docs/api/run-tests/reruntestexecutions","cat":"Docs"},{"title":"Delete Test Executions","desc":"Bulk-delete test executions from a test run. Specify execution UUIDs or use selectAll. Active executions (running/pending/cancelling) cannot be deleted.","href":"https://docs.futureagi.com/docs/api/run-tests/deletetestexecutions","cat":"Docs"},{"title":"Get Execution Details","desc":"Retrieve a test execution by UUID with paginated call executions. Supports search, filters, and row grouping. Returns status, call counts, eval scores, and transcripts.","href":"https://docs.futureagi.com/docs/api/test-executions/gettestexecutiondetails","cat":"Docs"},{"title":"Get Execution KPIs","desc":"Retrieve KPI metrics for a test execution. Returns avg score, response time, call connection rate, latency, WPM, interruption rates, and per-eval averages.","href":"https://docs.futureagi.com/docs/api/test-executions/getkpis","cat":"Docs"},{"title":"Get Performance Summary","desc":"Retrieve pass/fail rates and top-performing scenarios for a test execution. Returns overall pass rate, total test runs, latest fail rate, and top 4 scenarios by score.","href":"https://docs.futureagi.com/docs/api/test-executions/getperformancesummary","cat":"Docs"},{"title":"Cancel Execution","desc":"Cancel an in-progress test execution by UUID. Execution must be in pending, running, or evaluating state. Returns success flag and cancellation confirmation message.","href":"https://docs.futureagi.com/docs/api/test-executions/cancelexecution","cat":"Docs"},{"title":"Rerun Calls","desc":"Rerun call executions within a test execution. Choose eval_only to re-evaluate existing data or call_and_eval to re-run calls. Returns success and failure counts per call.","href":"https://docs.futureagi.com/docs/api/test-executions/reruncalls","cat":"Docs"},{"title":"Get Call Details","desc":"Retrieve details of a specific call execution by UUID, including provider call ID, metrics, and full execution data.","href":"https://docs.futureagi.com/docs/api/test-executions/getcallexecutiondetails","cat":"Docs"},{"title":"Get Simulation Metrics","desc":"Retrieve latency, cost, and conversation metrics for a simulation run. Query by run test name, execution ID, or call execution ID. Supports pagination for run-level results.","href":"https://docs.futureagi.com/docs/api/simulation-analytics/metrics","cat":"Docs"},{"title":"Get Simulation Runs","desc":"Retrieve simulation run records with eval scores, scenario metadata, and per-call breakdowns. Query by run test name, execution ID, or call execution ID. Supports FMA summary.","href":"https://docs.futureagi.com/docs/api/simulation-analytics/runs","cat":"Docs"},{"title":"Get Simulation Analytics","desc":"Retrieve aggregated eval scores, per-metric averages, system summary, and Fix My Agent suggestions for a simulation run by run test name or execution ID.","href":"https://docs.futureagi.com/docs/api/simulation-analytics/analytics","cat":"Docs"},{"title":"List Datasets","desc":"Retrieve a paginated list of datasets in your organization. Supports page, page_size, name search, and sort query parameters.","href":"https://docs.futureagi.com/docs/api/datasets/list-datasets","cat":"Docs"},{"title":"Create Dataset","desc":"Create a new dataset with a specified name, number of rows, and number of columns. Returns the created dataset ID and metadata.","href":"https://docs.futureagi.com/docs/api/datasets/create-dataset","cat":"Docs"},{"title":"Create Empty Dataset","desc":"Create a new empty dataset with a given name, model type, and optional pre-created blank rows. Returns the new dataset ID.","href":"https://docs.futureagi.com/docs/api/datasets/create-empty-dataset","cat":"Docs"},{"title":"Upload Dataset from File","desc":"Create a new dataset by uploading a local CSV or JSONL file as multipart form data. Returns the new dataset ID and row count.","href":"https://docs.futureagi.com/docs/api/datasets/upload-dataset","cat":"Docs"},{"title":"Create from HuggingFace","desc":"Create a new dataset by importing rows from a HuggingFace dataset. Specify dataset name, config, split, and number of rows.","href":"https://docs.futureagi.com/docs/api/datasets/create-dataset-from-huggingface","cat":"Docs"},{"title":"Clone Dataset","desc":"Create a full copy of an existing dataset, preserving all columns and rows, under a new name specified in the request body.","href":"https://docs.futureagi.com/docs/api/datasets/clone-dataset","cat":"Docs"},{"title":"Duplicate Dataset","desc":"Create a new dataset from selected rows of an existing dataset. Supports duplicating all rows or a specific subset by row IDs.","href":"https://docs.futureagi.com/docs/api/datasets/duplicate-dataset","cat":"Docs"},{"title":"Add as New Dataset","desc":"Create a new dataset from selected columns of an existing dataset or experiment, with a custom column mapping and dataset name.","href":"https://docs.futureagi.com/docs/api/datasets/add-as-new","cat":"Docs"},{"title":"Update Dataset","desc":"Update properties of an existing dataset such as its name. Pass the dataset UUID as a path parameter and the new values in the body.","href":"https://docs.futureagi.com/docs/api/datasets/update-dataset","cat":"Docs"},{"title":"Merge Dataset","desc":"Merge rows from a source dataset into a target dataset. Choose to merge all rows or a selected subset using the request body.","href":"https://docs.futureagi.com/docs/api/datasets/merge-dataset","cat":"Docs"},{"title":"Delete Dataset","desc":"Permanently delete one or more datasets by passing an array of dataset UUIDs. Returns a count of successfully deleted datasets.","href":"https://docs.futureagi.com/docs/api/datasets/delete-dataset","cat":"Docs"},{"title":"Add Rows from File","desc":"Append rows to an existing dataset by uploading a CSV or JSONL file as multipart form data with the target dataset ID.","href":"https://docs.futureagi.com/docs/api/datasets/add-rows-from-file","cat":"Docs"},{"title":"Add Empty Rows","desc":"Append a specified number of empty rows to an existing dataset by ID. Returns a confirmation message on success.","href":"https://docs.futureagi.com/docs/api/datasets/add-empty-rows","cat":"Docs"},{"title":"Add Rows from Existing","desc":"Copy rows from a source dataset into a target dataset using a column mapping to align fields between the two datasets.","href":"https://docs.futureagi.com/docs/api/datasets/add-rows-from-existing","cat":"Docs"},{"title":"Add Rows from HuggingFace","desc":"Import rows from a HuggingFace dataset into an existing dataset. Specify dataset name, config, split, and number of rows to import.","href":"https://docs.futureagi.com/docs/api/datasets/add-rows-from-huggingface","cat":"Docs"},{"title":"Duplicate Rows","desc":"Create one or more copies of specific rows within a dataset by providing row UUIDs and the number of copies to produce.","href":"https://docs.futureagi.com/docs/api/datasets/duplicate-rows","cat":"Docs"},{"title":"Delete Rows","desc":"Delete one or more rows from a dataset by row UUIDs, or delete all rows at once using the selected_all_rows flag.","href":"https://docs.futureagi.com/docs/api/datasets/delete-rows","cat":"Docs"},{"title":"Update Cell Value","desc":"Update the value of a specific cell in a dataset by providing the dataset ID, row ID, column ID, and the new cell value.","href":"https://docs.futureagi.com/docs/api/datasets/update-cell-value","cat":"Docs"},{"title":"Get Column Details","desc":"Retrieve column metadata for a dataset including names, types, and IDs. Filter by source type or include run prompt columns.","href":"https://docs.futureagi.com/docs/api/datasets/columns/get-column-details","cat":"Docs"},{"title":"Get Column Config","desc":"Retrieve configuration details for a specific dataset column by UUID, including model, messages, and source-specific settings.","href":"https://docs.futureagi.com/docs/api/datasets/columns/get-column-config","cat":"Docs"},{"title":"Add Static Column","desc":"Add a single static column to an existing dataset by specifying the column name and data type. Returns the new column's metadata.","href":"https://docs.futureagi.com/docs/api/datasets/columns/add-static-column","cat":"Docs"},{"title":"Add Multiple Static Columns","desc":"Add multiple static columns to a dataset in one request. Each column requires a name and data type such as text or number.","href":"https://docs.futureagi.com/docs/api/datasets/columns/add-multiple-static-columns","cat":"Docs"},{"title":"Add Columns","desc":"Add one or more columns to a dataset in a single request. Specify each column's name and data type in the request body.","href":"https://docs.futureagi.com/docs/api/datasets/columns/add-columns","cat":"Docs"},{"title":"Update Column Name","desc":"Rename an existing column in a dataset by providing the new column name. Requires dataset UUID and column UUID as path params.","href":"https://docs.futureagi.com/docs/api/datasets/columns/update-column-name","cat":"Docs"},{"title":"Update Column Type","desc":"Change the data type of an existing dataset column. Use the preview flag to dry-run the conversion before committing the change.","href":"https://docs.futureagi.com/docs/api/datasets/columns/update-column-type","cat":"Docs"},{"title":"Delete Column","desc":"Permanently delete a column and all its cell data from a dataset. Requires both the dataset UUID and column UUID as path params.","href":"https://docs.futureagi.com/docs/api/datasets/columns/delete-column","cat":"Docs"},{"title":"Add Run Prompt Column","desc":"Add a run prompt column to a dataset that generates LLM responses per row. Configure model, messages, and output format in the body.","href":"https://docs.futureagi.com/docs/api/datasets/run-prompt/add-run-prompt-column","cat":"Docs"},{"title":"Edit Run Prompt Column","desc":"Update the model, messages, or output config of an existing run prompt column and re-execute it across all dataset rows.","href":"https://docs.futureagi.com/docs/api/datasets/run-prompt/edit-run-prompt-column","cat":"Docs"},{"title":"Get Run Prompt Config","desc":"Retrieve the full configuration of an existing run prompt column by column UUID, including model, messages, and output settings.","href":"https://docs.futureagi.com/docs/api/datasets/run-prompt/retrieve-run-prompt-column-config","cat":"Docs"},{"title":"Get Run Prompt Options","desc":"Retrieve available models, tools, output formats, and tool choices for configuring a run prompt column in a dataset.","href":"https://docs.futureagi.com/docs/api/datasets/run-prompt/retrieve-run-prompt-options","cat":"Docs"},{"title":"Get Model Voices","desc":"Retrieve the available voice options for a specific model's audio output. Pass the model name as a required query parameter.","href":"https://docs.futureagi.com/docs/api/datasets/run-prompt/get-model-voices","cat":"Docs"},{"title":"TTS Voices","desc":"List custom text-to-speech voices configured for your organization, including voice ID, provider, and display name for each entry.","href":"https://docs.futureagi.com/docs/api/datasets/run-prompt/tts-voices","cat":"Docs"},{"title":"Get Column Values","desc":"Retrieve sample cell values from specified dataset columns, mapped by placeholder name, for use in prompt preview and testing.","href":"https://docs.futureagi.com/docs/api/datasets/run-prompt/get-column-values","cat":"Docs"},{"title":"Run Prompt Stats","desc":"Get aggregated statistics for run prompt columns in a dataset, including average token usage, cost, and per-prompt-ID breakdowns.","href":"https://docs.futureagi.com/docs/api/datasets/analytics/run-prompt-stats","cat":"Docs"},{"title":"Eval Stats","desc":"Retrieve evaluation statistics for a dataset by column, including template name, metric count, and average score per eval column.","href":"https://docs.futureagi.com/docs/api/datasets/analytics/eval-stats","cat":"Docs"},{"title":"Annotation Summary","desc":"Get annotation statistics for a dataset, including total annotations, label distributions, and per-annotator contribution breakdown.","href":"https://docs.futureagi.com/docs/api/datasets/analytics/annotation-summary","cat":"Docs"},{"title":"Explanation Summary","desc":"Get an AI-generated natural language summary of a dataset's content and patterns, returned as a structured explanation object.","href":"https://docs.futureagi.com/docs/api/datasets/analytics/explanation-summary","cat":"Docs"},{"title":"Create Score","desc":"Create a single annotation score on a trace, span, or session. Accepts source type, source UUID, label UUID, score value, and score source type.","href":"https://docs.futureagi.com/docs/api/annotations/scores/create-score","cat":"Docs"},{"title":"Bulk Create Scores","desc":"Create multiple annotation scores on a single source in one request. Accepts source type, source UUID, and an array of label ID, value, and score source.","href":"https://docs.futureagi.com/docs/api/annotations/scores/bulk-create-scores","cat":"Docs"},{"title":"Get Scores for Source","desc":"Retrieve all annotation scores for a specific source (trace, span, generation, or session). Filter by label UUID or annotator UUID to narrow results.","href":"https://docs.futureagi.com/docs/api/annotations/scores/get-scores-for-source","cat":"Docs"},{"title":"List Scores","desc":"List annotation scores with optional filters for source type, source UUID, label, and annotator. Supports pagination via page and page_size parameters.","href":"https://docs.futureagi.com/docs/api/annotations/scores/list-scores","cat":"Docs"},{"title":"Delete Score","desc":"Soft-delete an annotation score by UUID. Only the score creator or an org admin can delete. Returns HTTP 204 on success.","href":"https://docs.futureagi.com/docs/api/annotations/scores/delete-score","cat":"Docs"},{"title":"Create Label","desc":"Create a new annotation label with a name, type (star, binary, categorical), description, settings, and optional notes. Returns the created label object.","href":"https://docs.futureagi.com/docs/api/annotations/labels/create-label","cat":"Docs"},{"title":"List Labels","desc":"List annotation labels with optional filters for type, name search, project, dataset scope, and usage count. Returns a paginated array of label objects.","href":"https://docs.futureagi.com/docs/api/annotations/labels/list-labels","cat":"Docs"},{"title":"Get Label","desc":"Retrieve a specific annotation label by UUID, including its type, settings, description, project scope, and allow-notes configuration.","href":"https://docs.futureagi.com/docs/api/annotations/labels/get-label","cat":"Docs"},{"title":"Update Label","desc":"Update an existing annotation label's name, description, type, or settings by UUID. Returns the updated label object with all current field values.","href":"https://docs.futureagi.com/docs/api/annotations/labels/update-label","cat":"Docs"},{"title":"Delete Label","desc":"Soft-delete an annotation label by UUID. Deleted labels are hidden but recoverable via the restore endpoint. Returns HTTP 204 on success.","href":"https://docs.futureagi.com/docs/api/annotations/labels/delete-label","cat":"Docs"},{"title":"Restore Label","desc":"Restore a previously soft-deleted annotation label by UUID. Returns the full restored label object with its original type, settings, and configuration.","href":"https://docs.futureagi.com/docs/api/annotations/labels/restore-label","cat":"Docs"},{"title":"Create Queue","desc":"Create an annotation queue with name, assignment strategy, annotator list, labels, and review settings. Returns the fully configured queue object.","href":"https://docs.futureagi.com/docs/api/annotations/queues/create-queue","cat":"Docs"},{"title":"List Queues","desc":"List annotation queues with optional filters for status and name search. Supports pagination and optionally includes item counts per queue.","href":"https://docs.futureagi.com/docs/api/annotations/queues/list-queues","cat":"Docs"},{"title":"Get Queue","desc":"Retrieve full details of an annotation queue by UUID, including assignment strategy, labels, annotators, review settings, and creation metadata.","href":"https://docs.futureagi.com/docs/api/annotations/queues/get-queue","cat":"Docs"},{"title":"Update Queue","desc":"Update an annotation queue's name, description, instructions, assignment strategy, or review settings by UUID. Returns the updated queue configuration.","href":"https://docs.futureagi.com/docs/api/annotations/queues/update-queue","cat":"Docs"},{"title":"Delete Queue","desc":"Soft-delete an annotation queue by UUID. The queue and its items are hidden but not permanently removed. Returns a deleted boolean confirmation.","href":"https://docs.futureagi.com/docs/api/annotations/queues/delete-queue","cat":"Docs"},{"title":"Update Status","desc":"Transition an annotation queue to a new status (draft, active, paused, or completed). Accepts queue UUID and target status. Returns the updated queue object.","href":"https://docs.futureagi.com/docs/api/annotations/queues/update-status","cat":"Docs"},{"title":"Get Progress","desc":"Retrieve progress statistics for an annotation queue, including total, pending, in-progress, completed, and skipped counts with per-annotator breakdowns.","href":"https://docs.futureagi.com/docs/api/annotations/queues/get-progress","cat":"Docs"},{"title":"Get Analytics","desc":"Retrieve detailed analytics for a queue including daily throughput, annotator performance, label score distribution, and item status breakdown.","href":"https://docs.futureagi.com/docs/api/annotations/queues/get-analytics","cat":"Docs"},{"title":"Get Agreement","desc":"Retrieve inter-annotator agreement scores for a queue, including overall agreement percentage and per-label breakdown for quality assurance analysis.","href":"https://docs.futureagi.com/docs/api/annotations/queues/get-agreement","cat":"Docs"},{"title":"Export","desc":"Export annotation queue items and their scores as JSON or CSV. Filter by item status. Returns source type, annotations, annotator names, and timestamps.","href":"https://docs.futureagi.com/docs/api/annotations/queues/export","cat":"Docs"},{"title":"Export to Dataset","desc":"Export completed annotation queue items into a Future AGI dataset. Accepts queue UUID, target dataset ID, and optional status filter for selective export.","href":"https://docs.futureagi.com/docs/api/annotations/queues/export-to-dataset","cat":"Docs"},{"title":"Add Label to Queue","desc":"Attach an existing annotation label to a queue by queue UUID and label UUID. Returns the updated queue-label association with ordering and required flag.","href":"https://docs.futureagi.com/docs/api/annotations/queues/add-label","cat":"Docs"},{"title":"Remove Label","desc":"Detach an annotation label from a queue by queue UUID and label UUID. Removing a label hides it from future annotation sessions for that queue.","href":"https://docs.futureagi.com/docs/api/annotations/queues/remove-label","cat":"Docs"},{"title":"Get or Create Default","desc":"Get the default annotation queue for a project, dataset, or agent definition, automatically creating one if it does not yet exist. Returns queue and labels.","href":"https://docs.futureagi.com/docs/api/annotations/queues/get-or-create-default","cat":"Docs"},{"title":"Find Queues for Source","desc":"Find all annotation queues containing a specific source item. Accepts source type, source UUID, or a JSON array of sources for multi-source lookup.","href":"https://docs.futureagi.com/docs/api/annotations/queues/find-queues-for-source","cat":"Docs"},{"title":"List Items","desc":"List items in an annotation queue with pagination. Filter by status, source type, or assigned user. Returns item details and annotation progress.","href":"https://docs.futureagi.com/docs/api/annotations/items/list-items","cat":"Docs"},{"title":"Add Items","desc":"Add one or more source items (traces, spans) to an annotation queue in bulk. Accepts queue UUID and an array of source type and source ID pairs.","href":"https://docs.futureagi.com/docs/api/annotations/items/add-items","cat":"Docs"},{"title":"Bulk Remove Items","desc":"Remove multiple items from an annotation queue in a single request. Accepts the queue UUID and an array of item UUIDs to delete from the queue.","href":"https://docs.futureagi.com/docs/api/annotations/items/bulk-remove-items","cat":"Docs"},{"title":"Get Annotate Detail","desc":"Retrieve a queue item with full source content, queue metadata, labels, existing annotations, and navigation pointers for the annotation interface.","href":"https://docs.futureagi.com/docs/api/annotations/items/get-annotate-detail","cat":"Docs"},{"title":"Get Next Item","desc":"Retrieve the next available queue item for the current user to annotate. Returns item ID, source type, status, and assignment info.","href":"https://docs.futureagi.com/docs/api/annotations/items/get-next-item","cat":"Docs"},{"title":"Submit Annotations","desc":"Submit label scores and reviewer notes for an annotation queue item. Accepts queue UUID, item UUID, an array of annotation values, and optional notes.","href":"https://docs.futureagi.com/docs/api/annotations/items/submit-annotations","cat":"Docs"},{"title":"Complete Item","desc":"Mark an annotation queue item as completed. Returns the completed item ID and, optionally, the next pending item available for annotation.","href":"https://docs.futureagi.com/docs/api/annotations/items/complete-item","cat":"Docs"},{"title":"Skip Item","desc":"Mark an annotation queue item as skipped by the current user. Returns the skipped item ID and the next available pending item for annotation.","href":"https://docs.futureagi.com/docs/api/annotations/items/skip-item","cat":"Docs"},{"title":"Get Item Annotations","desc":"Retrieve all annotation scores submitted for a specific queue item, including label values, score source, annotator ID, and timestamps.","href":"https://docs.futureagi.com/docs/api/annotations/items/get-item-annotations","cat":"Docs"},{"title":"Assign Items","desc":"Assign one or more annotation queue items to specific annotators. Accepts queue UUID, a list of item UUIDs, and a list of user UUIDs to assign.","href":"https://docs.futureagi.com/docs/api/annotations/items/assign-items","cat":"Docs"},{"title":"Release Item","desc":"Release a reserved annotation queue item so it becomes available for reassignment to another annotator. Returns a boolean release confirmation.","href":"https://docs.futureagi.com/docs/api/annotations/items/release-item","cat":"Docs"},{"title":"Bulk Annotate Spans","desc":"Submit label scores and notes for multiple observation spans in one request. Accepts an array of span IDs with their annotation values and score sources.","href":"https://docs.futureagi.com/docs/api/annotations/bulk/bulk-annotate-spans","cat":"Docs"},{"title":"Best 5 Error Analysis Tools for HR AI Agents in 2026","desc":"Five error analysis tools for HR AI agents in 2026: cluster HR failures, localize the root cause, write the fix, and redact candidate and employee PII at the span layer.","href":"/blog/error-analysis/error-analysis-tools-hr-ai-agents-2026","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for Insurance AI Agents in 2026","desc":"Five error analysis tools for insurance AI agents in 2026: cluster claims and underwriting failures, localize the root cause, write the fix, and redact PII at the span layer.","href":"/blog/error-analysis/error-analysis-tools-insurance-ai-agents-2026","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for Hospitality AI Agents in 2026","desc":"Five error analysis tools for hospitality AI agents in 2026: cluster booking and rate failures, localize the root cause, write the fix, and redact guest PII at the span layer.","href":"/blog/error-analysis/error-analysis-tools-hospitality-ai-agents-2026","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for Hospitality in 2026","desc":"The 5 best AI agent simulation tools for hospitality in 2026, scored on booking accuracy, multilingual guest handling, auto-scenario generation, and eval-linked verdicts. FutureAGI leads.","href":"/blog/ai-agent-simulation-tools-hospitality","cat":"Blog"},{"title":"How to Evaluate GraphRAG Pipelines: Metrics Beyond RAG","desc":"Base RAG metrics miss the graph underneath GraphRAG. Here is a three-layer framework, runnable graph metrics, and answer scores that isolate each failure.","href":"/blog/how-to-evaluate-graphrag-pipelines","cat":"Blog"},{"title":"LLM Regression Testing: Catch Silent Model-Swap Failures","desc":"LLM regression testing catches the silent model-swap failure that lowers quality with no error. Here is a runnable harness, a CI gate, and golden-set checks.","href":"/blog/llm-regression-testing-model-swap","cat":"Blog"},{"title":"Reward Model Drift in LLMs: How to Detect It","desc":"A reward model can decay silently while its scores keep climbing. Here are the exact detectors, thresholds, and pipeline placement to catch the drift.","href":"/blog/reward-model-drift-llm","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for Retail AI Agents in 2026","desc":"Five error analysis tools for retail AI agents in 2026: cluster retail failures, localize the root cause, write the fix, and redact customer PII at the span layer.","href":"/blog/error-analysis/error-analysis-tools-retail-ai-agents-2026","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for Legal in 2026","desc":"The 5 best AI agent simulation tools for legal teams in 2026, scored on citation integrity, scenario realism, and eval-linked verdicts.","href":"/blog/ai-agent-simulation-tools-legal","cat":"Blog"},{"title":"Intelligence Explosion vs Recursive Self-Improvement","desc":"Intelligence explosion vs recursive self-improvement: one is a loop running today, the other a hypothesized runaway outcome. Which claims belong to each.","href":"/blog/intelligence-explosion-vs-recursive-self-improvement","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for CX AI Agents in 2026","desc":"Five error analysis tools for CX AI agents in 2026: cluster support failures, localize the root cause, write the fix, and redact customer PII at the span layer.","href":"/blog/error-analysis/error-analysis-tools-cx-ai-agents-2026","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for HR in 2026","desc":"The 5 best AI agent simulation tools for HR in 2026, scored on EEOC bias coverage, scenario realism, and eval-linked verdicts.","href":"/blog/ai-agent-simulation-tools-hr","cat":"Blog"},{"title":"Recursive Self-Improvement in AI: 2026 Examples","desc":"How recursive self-improvement works in AI, the verified systems running the loop right now, and the fixed evaluation signal that keeps each one bounded.","href":"/blog/recursive-self-improvement-ai-2026-examples","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for Education AI Agents in 2026","desc":"Five error analysis tools for education AI agents in 2026: cluster education failures, localize the root cause, write the fix, and redact student PII at the span layer.","href":"/blog/error-analysis/error-analysis-tools-education-ai-agents-2026","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for CX in 2026","desc":"The 5 leading AI agent simulation tools for CX in 2026, scored on escalation handling, multi-turn CSAT, scenario realism, and eval-linked verdicts, with FutureAGI ranked first.","href":"/blog/ai-agent-simulation-tools-cx","cat":"Blog"},{"title":"AI Code Review Tools: Scoring Precision, Recall and False Flags","desc":"A method to score AI code review tools yourself: what precision, recall, and false discovery rate mean on a pull request, and the metric vendors report wrong.","href":"/blog/ai-code-review-tools-precision-recall","cat":"Blog"},{"title":"Text to SQL LLM Systems: How to Evaluate Generated Queries Before They Run","desc":"A query can execute cleanly and still return the wrong number. Here is a failure-mode taxonomy and a validation layer that runs before the database does.","href":"/blog/text-to-sql-llm-evaluation","cat":"Blog"},{"title":"AI Testing Tools for LLM Products: What Traditional QA Suites Cannot Check","desc":"Traditional test suites pass green while a model ships a wrong answer. See what AI testing tools check that assertions cannot, and how to gate it in CI.","href":"/blog/ai-testing-tools-llm-products","cat":"Blog"},{"title":"AI Audit Checklist: The Evaluation Evidence Your Stack Should Already Produce","desc":"Most audit prep is a scramble to reconstruct what happened. Map each requirement to a system that already records it, then automate the two or three that are left.","href":"/blog/ai-audit-checklist","cat":"Blog"},{"title":"AI Agent Security Risks in 2026: Testing Against Prompt Injection and Tool Abuse","desc":"Four disclosed incidents, most of them this year, turned agent security from a thought experiment into an engineering problem. Here is what to test for.","href":"/blog/ai-agent-security-risks","cat":"Blog"},{"title":"Human in the Loop AI: Designing Review Queues That Improve LLM Quality","desc":"Design a human in the loop AI review queue that lifts LLM quality: what to route, how to score reviewer agreement with kappa, and how to prove the queue worked.","href":"/blog/human-in-the-loop-ai-review-queues","cat":"Blog"},{"title":"HumanEval Benchmark Explained: What Pass@k Does and Does Not Prove","desc":"How HumanEval scores 164 problems, how the pass@k estimator works, and the contamination, weak-test, and single-function limits behind a high score.","href":"/blog/humaneval-benchmark","cat":"Blog"},{"title":"OpenAI API Rate Limits: Testing Provider Reliability Under Pressure","desc":"What RPM, TPM and the six spend tiers actually enforce, why 429s arrive under your cap, and how to check a provider holds up before launch day does it for you.","href":"/blog/openai-api-rate-limits","cat":"Blog"},{"title":"Perplexity AI Review: Where Its Citations Hold Up and Where They Break","desc":"A citation focused Perplexity AI review: what independent studies found about its sourcing, and a one minute check to run before you quote a linked source.","href":"/blog/perplexity-ai-review-citations","cat":"Blog"},{"title":"What Is Named Entity Recognition: How to Measure NER Quality in LLM Pipelines","desc":"Extraction that looks right can still be wrong. Here is how precision, recall, F1, and partial matching work when an LLM is doing the tagging.","href":"/blog/what-is-named-entity-recognition","cat":"Blog"},{"title":"What Is Prompt Chaining: How to Trace and Score Every Step","desc":"A chain that returns a clean answer can still be wrong in the middle. Here is how to record each link as its own span and put a score on it.","href":"/blog/what-is-prompt-chaining","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for Legal AI Agents in 2026","desc":"Five error analysis tools for legal AI agents in 2026: cluster legal failures, localize the root cause, write the fix, and redact privileged data at the span layer.","href":"/blog/error-analysis/error-analysis-tools-legal-ai-agents-2026","cat":"Blog"},{"title":"AI Model Governance: Turning Model Choice Into a Repeatable Approval Process","desc":"A four-step approval sequence, a model scorecard template, and a maturity checklist you can copy to make every model decision documented and reviewable.","href":"/blog/ai-model-governance","cat":"Blog"},{"title":"Computer Use Agent Evaluation: How Screen-Acting Agents Fail and How to Score Them","desc":"A plain-language failure taxonomy for screen-acting agents, plus a five-dimension scoring rubric that goes past pass/fail task completion.","href":"/blog/computer-use-agent-evaluation","cat":"Blog"},{"title":"LLM Benchmarks Explained: A Methodology Checklist Before You Trust a Score","desc":"How benchmark scores get inflated by contamination, selective reporting, and sampling noise, plus a six-criterion checklist for reading any leaderboard.","href":"/blog/llm-benchmarks-explained","cat":"Blog"},{"title":"What Is Temperature in LLM Output: What to Measure Before You Tune","desc":"How temperature reshapes token probabilities, why temperature 0 still drifts, and the variance and faithfulness checks to run before locking a value.","href":"/blog/llm-temperature","cat":"Blog"},{"title":"What Is a State Machine: Deterministic Control for Multi Step Agent Flows","desc":"States, transitions, and guards with a worked transition table, plus how the pattern bounds multi-step agent flows and how to check the path an agent took.","href":"/blog/state-machine","cat":"Blog"},{"title":"Data Lineage Tools for AI Systems: Proving Which Source Produced an Answer","desc":"Data lineage tools trace data across your stack, but the graph stops at the model boundary. Here are the two records that prove which source shaped an answer.","href":"/blog/data-lineage-tools-ai-systems","cat":"Blog"},{"title":"Why Did My RAG Agent Get Worse? A Live Autopsy","desc":"A live RAG agent autopsy with Future AGI and Qdrant: trace a RAG pipeline, find why retrieval degrades as data grows, and fix it from 52% to 92% accuracy.","href":"/blog/why-did-my-rag-agent-get-worse-webinar-2026","cat":"Blog"},{"title":"Observability vs Monitoring: What the Difference Means for AI Agents","desc":"The observability vs monitoring difference decides whether you catch a wrong-but-successful agent answer that green dashboards hide, and why evals catch it.","href":"/blog/observability-vs-monitoring","cat":"Blog"},{"title":"OpenTelemetry vs Prometheus for LLM Telemetry: Traces and Metrics","desc":"The opentelemetry vs prometheus question for LLM apps, answered: what each tool does, why they are complementary, the gen_ai conventions, and how to run both.","href":"/blog/opentelemetry-vs-prometheus","cat":"Blog"},{"title":"When to Move From an Agentic Loop to a Simpler Deterministic Workflow","desc":"A plain decision framework for when to drop an agentic loop for a simpler deterministic workflow: the three questions to ask, the signals to watch, and how to convert one.","href":"/blog/loop-engineering/agentic-loop-vs-deterministic-workflow","cat":"Blog"},{"title":"Self-Correcting Agent Loops: How Agents Detect and Fix Their Own Mistakes","desc":"How a self-correcting agent loop spots its own bad output and revises it: the generate, critique, revise structure, the patterns behind it, and when it backfires.","href":"/blog/loop-engineering/self-correcting-agent-loops","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for Retail in 2026","desc":"The 5 leading AI agent simulation tools for retail in 2026, scored on scenario realism, auto-scenario generation, and eval-linked verdicts for brand-voice drift, PDP pricing accuracy, and returns.","href":"/blog/ai-agent-simulation-tools-retail","cat":"Blog"},{"title":"Best Open-Source AI Agent Simulation Tools in 2026","desc":"The best open-source AI agent simulation tools in 2026, scored on license, self-hosting, and openness. FutureAGI ranks first as the only fully Apache 2.0 simulate to evaluate to observe loop.","href":"/blog/open-source-ai-agent-simulation-tools","cat":"Blog"},{"title":"MCP vs API: How Model Context Protocol Changes Agent Tool Access","desc":"The mcp vs api question for agent builders: how Model Context Protocol changes tool access, how it compares to REST and function calling, and when to use each.","href":"/blog/mcp-vs-api","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for Healthcare AI Agents in 2026","desc":"Five error analysis tools for healthcare AI agents in 2026: cluster clinical failures, localize the root cause, write the fix, and redact PHI at the span layer.","href":"/blog/error-analysis/error-analysis-tools-healthcare-ai-agents-2026","cat":"Blog"},{"title":"Best LLMs of July 2026: Top Closed-Source, Open-Weight, Multimodal, and Coding Picks","desc":"Best LLMs of July 2026 by use case: Claude Opus 5 for agentic coding, GPT-5.6 Sol for reasoning, Kimi K3 for open-weight scale, Gemini 3.6 Flash for speed.","href":"/blog/best-llms-july-2026","cat":"Blog"},{"title":"Best Voice AI Models in July 2026: STT, TTS, and Voice Agent Stack","desc":"STT, TTS, and voice-agent picks for July 2026: ElevenLabs Scribe v2 leads accuracy at 2.2% WER, Deepgram Nova-3 for streaming, Cartesia Sonic-3.5 for TTS.","href":"/blog/best-voice-ai-july-2026","cat":"Blog"},{"title":"How to Stop an AI Agent Loop From Burning Through Your Budget","desc":"How an AI agent infinite loop runs up token cost, and the pre-call budget, guards, and cost levers that stop the bleed before it hits your invoice.","href":"/blog/loop-engineering/ai-agent-loop-cost-control","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools in 2026: Ranked for Production","desc":"The 5 best AI agent simulation tools in 2026, ranked for production reliability: scale, drift, and live monitoring.","href":"/blog/ai-agent-simulation-tools-production","cat":"Blog"},{"title":"Future AGI Q2 2026: Fully Open Source Under Apache 2.0","desc":"Inside Future AGI open source in Q2 2026: the platform shipped under Apache 2.0, Error Feed and the Agent Command Center went live, traces hit billions.","href":"/blog/future-agi-q2-2026-open-source","cat":"Blog"},{"title":"Best 5 Error Analysis Tools for Fintech AI Agents in 2026","desc":"Five error analysis tools for fintech AI agents in 2026: cluster compliance failures, localize the root cause, write the fix, and redact PII at the span layer.","href":"/blog/error-analysis/error-analysis-tools-fintech-ai-agents-2026","cat":"Blog"},{"title":"The Agentic Loop Explained: Lifecycle, Diagram, and Examples","desc":"A plain walkthrough of the agentic loop: the named lifecycle stages, the canonical diagram, a worked example, and what decides when the loop stops.","href":"/blog/loop-engineering/agentic-loop-explained","cat":"Blog"},{"title":"Why Your AI Agent Gets Stuck in a Loop (and How to Fix It)","desc":"Why AI agents get stuck repeating the same action, the three root causes behind it, and the fixes that actually stop the loop instead of just delaying it.","href":"/blog/loop-engineering/ai-agent-stuck-in-loop","cat":"Blog"},{"title":"Designing Agentic Loops: Sandboxing, Safety, and Letting Agents Run Code","desc":"How to let an agent run code without wrecking your machine: what a sandbox boundary limits, how to scope filesystem and network access, and where guardrails fit.","href":"/blog/loop-engineering/designing-agentic-loops","cat":"Blog"},{"title":"The Ralph Loop: Running a Coding Agent in an Autonomous Loop","desc":"How the Ralph loop works: a shell loop that re-runs a coding agent with fresh context each pass, the tools that add a stop condition, and what it costs.","href":"/blog/loop-engineering/ralph-loop","cat":"Blog"},{"title":"ReAct Agent Loop: How Reason-and-Act Agents Actually Work","desc":"How a ReAct agent loop works: the Thought, Action, Observation cycle, the paper it came from, and why interleaving reasoning with acting beats planning everything up front.","href":"/blog/loop-engineering/react-agent-loop","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for Education in 2026","desc":"The 5 leading AI agent simulation tools for education in 2026, scored on FERPA and minor-safety coverage, multi-turn academic realism, and eval-linked verdicts. FutureAGI ranks first.","href":"/blog/ai-agent-simulation-tools-education","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for Fintech in 2026","desc":"The 5 best AI agent simulation tools for fintech in 2026, scored on SR 11-7 fit, scenario realism, and eval-linked verdicts. FutureAGI ranks first.","href":"/blog/ai-agent-simulation-tools-fintech","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for Insurance in 2026","desc":"The 5 leading AI agent simulation tools for insurance in 2026, scored on the 5-Criteria Simulation Scorecard: scenario realism, auto-scenario generation, eval-linked verdicts, DOI and NAIC bias coverage, and deployment.","href":"/blog/ai-agent-simulation-tools-insurance","cat":"Blog"},{"title":"Prompt, Context, Harness, Loop: The Four Layers of AI Agent Engineering","desc":"Prompt, context, harness, and loop engineering explained as four layers of one agent stack, including where harness engineering vs context engineering actually differ.","href":"/blog/loop-engineering/prompt-context-harness-loop-layers","cat":"Blog"},{"title":"How to Actually Do Loop Engineering: Worktrees, Tests, and Cost Caps","desc":"A practitioner's guide to building agent loops that hold up under real use: isolating parallel runs with git worktrees, gating changes on a real test command, and capping cost before you scale.","href":"/blog/loop-engineering/how-to-do-loop-engineering","cat":"Blog"},{"title":"Top 5 AI Agent Error Analysis Tools in 2026: Ranked for Production","desc":"Five AI agent error analysis tools for ML and platform teams in 2026: failure clustering, root-cause attribution, drift, and the fix loop that closes the gap.","href":"/blog/error-analysis/ai-agent-error-analysis-tools-2026","cat":"Blog"},{"title":"Is Loop Engineering Real, or Just Another AI Buzzword?","desc":"An honest look at the loop engineering debate: the real case that it's a genuinely new practice, the real case that it's renamed orchestration, and a clear verdict.","href":"/blog/loop-engineering/is-loop-engineering-real","cat":"Blog"},{"title":"Top 5 AI Agent Simulation Tools for Healthcare in 2026","desc":"The 5 leading AI agent simulation tools for healthcare in 2026, scored on HIPAA, PHI, multi-turn realism, and eval-linked verdicts. FutureAGI ranks first.","href":"/blog/ai-agent-simulation-tools-healthcare","cat":"Blog"},{"title":"How to Simulate AI Agents with Open-Source Tools in 2026","desc":"A step-by-step walkthrough for how to simulate AI agents with open-source tools: install the SDK, write personas and scenarios, run the test, and fix failures before you ship.","href":"/blog/how-to-simulate-ai-agents-open-source","cat":"Blog"},{"title":"What Is an AI Agent Loop? The Core Architecture Behind Autonomous Agents","desc":"A plain breakdown of what actually sits inside an AI agent loop: the model, the memory, the planner, the tools, and the controller that ties them together.","href":"/blog/loop-engineering/what-is-ai-agent-loop","cat":"Blog"},{"title":"Loop Engineering vs Prompt Engineering: Why the Shift Matters","desc":"Prompt engineering shapes one request. Loop engineering designs the system that decides what to request next. Here is what actually changed, and what didn't.","href":"/blog/loop-engineering/loop-engineering-vs-prompt-engineering","cat":"Blog"},{"title":"Loop Engineering vs Harness Engineering: What's the Difference?","desc":"Loop engineering and harness engineering control different layers of the same agent system. Here is how to tell which one you are actually missing.","href":"/blog/loop-engineering/loop-engineering-vs-harness-engineering","cat":"Blog"},{"title":"AI Agent Simulation in 2026: A Practical Guide","desc":"A practical guide to AI agent simulation: what it is, the 3-layer stack of personas, scenarios, and eval-linked verdicts, and how to simulate an agent.","href":"/blog/ai-agent-simulation-practical-guide-2026","cat":"Blog"},{"title":"Open-Source AI News: July 2026 LLM Release Tracker","desc":"A verified July 2026 open-source LLM release tracker separating announcements from downloadable weights, original models from derivatives, and license types.","href":"/blog/open-source-llm-releases","cat":"Blog"},{"title":"What Is Loop Engineering? How Designing Agent Loops Replaces Prompting","desc":"A plain-language definition of loop engineering: who coined it, what a working loop actually contains, and when building one is worth the effort.","href":"/blog/loop-engineering/what-is-loop-engineering","cat":"Blog"},{"title":"Agent Harness vs Framework: What Is the Difference?","desc":"Agent harness vs framework: a framework supplies reusable abstractions, while a harness runs and controls the agent loop. Learn when you need each.","href":"/blog/agent-harness-vs-framework","cat":"Blog"},{"title":"Best Open-Source AI Projects to Contribute to in 2026","desc":"Compare eight open-source AI projects using live GitHub activity, newcomer issues, contribution docs, licensing, and merge throughput.","href":"/blog/best-open-source-projects","cat":"Blog"},{"title":"Open-Source vs Open-Weight LLMs: The OSI Test (2026)","desc":"Open-source and open-weight LLMs are not the same. Use the OSI test to check model weights, code, data information, licenses, and deployment rights.","href":"/blog/open-source-vs-open-weight","cat":"Blog"},{"title":"What Is an Agent Harness? Components and Reliability","desc":"An agent harness is the runtime layer around an LLM that manages tools, context, state, execution, and safety. Learn how it shapes reliability.","href":"/blog/what-is-an-agent-harness","cat":"Blog"},{"title":"What Is OSINT? Why Verification Beats Search (2026)","desc":"What is OSINT? Learn the five-step verification workflow, legal boundaries, source checks, and how AI changes open-source intelligence in 2026.","href":"/blog/what-is-osint","cat":"Blog"},{"title":"Best AI Agent Testing & Simulation Tools in 2026","desc":"A plain-language guide to the best AI agent testing and simulation tools for 2026, scored for CI/CD: regression suites, pre-deploy gates, eval-linked verdicts.","href":"/blog/ai-agent-testing-simulation-tools-2026","cat":"Blog"},{"title":"Agent Harness Architecture: The 70% of Performance That Lives Outside the Model","desc":"How agent harness architecture shapes performance across five dimensions, from context and tools to safety and the loop that often outweighs a model upgrade.","href":"/blog/agent-harness-architecture","cat":"Blog"},{"title":"SWE-bench Evaluation Harness: How to Run It Without the Docker Errors","desc":"Fix the common SWE-bench harness Docker failures: the 120 GB disk trap, cache_level tradeoffs, ARM64 builds, mid-run hangs, and manifest-not-found errors.","href":"/blog/swe-bench-evaluation-harness","cat":"Blog"},{"title":"Agent Eval Harness: How to Evaluate AI Agents, Not Just Models","desc":"Evaluating an agent means scoring trajectories, tool calls, and environment outcomes. How an agent eval harness works, and the benchmarks that power it.","href":"/blog/agent-eval-harness","cat":"Blog"},{"title":"The Best Agent Harnesses for AI Builders in 2026","desc":"A field guide to the best agent harnesses in 2026: Claude Code, Codex CLI, OpenHands, SWE-agent, Aider, Cline and more, and how to pick one for your stack.","href":"/blog/best-agent-harness","cat":"Blog"},{"title":"Best Enterprise Prompt Management Platform for Fintech AI in 2026","desc":"Prompt management for fintech AI: Future AGI versions and governs each prompt in your own VPC, flags PII at the boundary, and eval-gates every promotion.","href":"/blog/best-enterprise-prompt-management-platform-fintech-2026","cat":"Blog"},{"title":"Best Prompt Management Platform for Legal AI in 2026","desc":"Prompt management for legal AI: Future AGI self-hosts the prompt registry in your perimeter, flags PII on inputs and outputs, and eval-gates every promotion.","href":"/blog/best-prompt-management-platform-legal-ai-2026","cat":"Blog"},{"title":"How to Add a Custom Task to lm-evaluation-harness","desc":"Write a custom lm-evaluation-harness task in YAML: the required fields, prompt templates, output types, metric config, and loading it without forking the repo.","href":"/blog/lm-evaluation-harness-custom-task","cat":"Blog"},{"title":"What Is Open-Source Software? Licenses, Models & Why It Matters (2026)","desc":"Open-source software lets you use, study, and change code under an OSI license. Here are the license families, what does not qualify, and why it matters.","href":"/blog/what-is-open-source-software-2026","cat":"Blog"},{"title":"Agent Harness vs Runtime vs Framework: The Terms, Untangled","desc":"Framework, runtime, or harness: what each agent layer actually covers, where reputable sources disagree, and how to tell which one your stack is missing.","href":"/blog/agent-harness-vs-runtime","cat":"Blog"},{"title":"Best Prompt Management Platform for RAG Applications in 2026","desc":"Prompt management for RAG: Future AGI versions retrieval and synthesis prompts, gates them on groundedness evaluators, and traces every answer.","href":"/blog/best-prompt-management-platform-rag-applications-2026","cat":"Blog"},{"title":"Best Enterprise Prompt Management Platform for Healthcare AI in 2026","desc":"Prompt management for healthcare AI: Future AGI self-hosts the registry so PHI stays in your perimeter, flags PII, and gates each prompt on clinical evaluators.","href":"/blog/best-enterprise-prompt-management-platform-healthcare-2026","cat":"Blog"},{"title":"Harness Engineering: The Emerging Playbook for Building AI Agents","desc":"Harness engineering is the discipline of building execution environments that make AI agents reliable: five principles, the ratchet rule, and what to measure.","href":"/blog/harness-engineering","cat":"Blog"},{"title":"Best Prompt Versioning Tool for CI/CD Pipelines in 2026","desc":"Prompt versioning for CI/CD with Future AGI: commit each version from the SDK, gate promotion on evaluators, fail the build on regression, promote by label.","href":"/blog/best-prompt-versioning-tool-cicd-pipelines-2026","cat":"Blog"},{"title":"Agent Harness vs MCP: How They Fit Together","desc":"MCP is a protocol for connecting agents to tools; an agent harness is the runtime that uses it. How the two layers divide work and when you need each.","href":"/blog/agent-harness-vs-mcp","cat":"Blog"},{"title":"Best Prompt Management Platform for Voice AI in 2026","desc":"Prompt management for voice AI: Future AGI versions each voice-agent prompt, simulates it against synthetic callers, and gates promotion on resolution and tone.","href":"/blog/best-prompt-management-platform-voice-ai-2026","cat":"Blog"},{"title":"Choosing an AI Evaluation Platform for Healthcare AI: 5 Questions to Ask","desc":"Five questions for choosing an AI evaluation platform for healthcare AI: clinical accuracy checks, no ground truth needed, custom medical criteria, and error localization.","href":"/blog/choosing-x-tools-y-questions/choosing-ai-evaluation-healthcare","cat":"Blog"},{"title":"Best LLMs of June 2026: Top Closed-Source, Open-Weight, Multimodal, and Coding Picks","desc":"Best LLMs of June 2026 by use case: Claude Fable 5 for raw coding, GLM-5.2 for open-weight value, GPT-5.5 for agents, Gemini 3.1 Pro for long-context multimodal.","href":"/blog/best-llms-june-2026","cat":"Blog"},{"title":"Best Prompt Versioning Tool for Production AI in 2026","desc":"Prompt versioning for production AI with Future AGI: immutable snapshots, each trace stamped with the version that ran, eval-gated promotion, fast rollback.","href":"/blog/best-prompt-versioning-tool-production-ai-2026","cat":"Blog"},{"title":"Best Voice AI Models in June 2026: STT, TTS, and Voice Agent Stack","desc":"Best Voice AI June 2026: Deepgram Nova-3 for streaming STT, Cartesia Sonic-3.5 for TTS, Retell for voice agents, plus latency budgets and cost at scale.","href":"/blog/best-voice-ai-june-2026","cat":"Blog"},{"title":"lm-evaluation-harness: A Practical Guide to EleutherAI's LLM Eval Tool","desc":"Install, run, and understand EleutherAI's lm-evaluation-harness: the CLI flags, the backends it supports, task structure, and the first-run pitfalls to avoid.","href":"/blog/lm-evaluation-harness-guide","cat":"Blog"},{"title":"Best Prompt Management Platform for Multi-Agent Systems in 2026","desc":"Prompt management for multi-agent systems: Future AGI versions, evaluates, and traces each agent's prompt on its own in one registry. Apache-2.0.","href":"/blog/best-prompt-management-platform-multi-agent-2026","cat":"Blog"},{"title":"Coding Agent Harness Benchmarks: Why the Harness Changes the Score","desc":"The same model scores differently across coding agent harness benchmarks. Here's what a leaderboard number really encodes and how to read scaffold effects.","href":"/blog/coding-agent-harness-benchmark","cat":"Blog"},{"title":"Choosing a Conversation Simulation Tool for Chat Agents: 5 Questions to Ask","desc":"Five questions for choosing a conversation simulation tool for chat agents: realistic multi-turn dialogue, user personas, automatic scoring, and scenario coverage.","href":"/blog/choosing-x-tools-y-questions/choosing-conversation-simulation-chat-agents","cat":"Blog"},{"title":"Best Prompt Management Platform for Customer Support AI in 2026","desc":"Prompt management for customer support AI: Future AGI versions each support prompt, gates promotion on Conversation Resolution and Tone, and traces every chat.","href":"/blog/best-prompt-management-platform-customer-support-2026","cat":"Blog"},{"title":"Best Prompt Management Platform for AI Agents in 2026","desc":"Version each agent prompt, gate promotion on task-completion and tool-call evaluators, then trace every run back to its version. Future AGI is Apache-2.0.","href":"/blog/best-prompt-management-platform-ai-agents-2026","cat":"Blog"},{"title":"Build a Minimal Agent Harness in Python (Step by Step)","desc":"A step-by-step Python guide to the agent harness loop: tool schema setup, stop_reason handling, parallel tool calls, and the bugs every builder hits first.","href":"/blog/build-agent-harness-python","cat":"Blog"},{"title":"Choosing an LLM Observability Tool for Your Stack: 7 Questions","desc":"Seven questions for choosing an LLM observability tool that fits your existing stack: OpenTelemetry-native spans, auto-instrumentation, portable export, and eval scoring.","href":"/blog/choosing-x-tools-y-questions/choosing-llm-observability-existing-stack","cat":"Blog"},{"title":"Best Enterprise Prompt Management Platforms in 2026","desc":"The 7 best enterprise prompt management platforms in 2026, compared on access control, versioning, evaluation, and security for large, regulated AI teams.","href":"/blog/best-enterprise-prompt-management-platforms-in-2026","cat":"Blog"},{"title":"What Is an Eval Harness? LLM Evaluation, Explained","desc":"An eval harness is the software that turns an LLM benchmark into a reproducible score. See how it loads tasks, formats prompts, scores outputs, and logs.","href":"/blog/what-is-an-eval-harness","cat":"Blog"},{"title":"Choosing AI Guardrails for Legal AI: 8 Factors to Compare","desc":"Eight factors for choosing an AI guardrails platform for legal AI: real-time blocking, confidentiality, fabricated-citation detection, tool permissioning, and self-hosting.","href":"/blog/choosing-x-tools-y-questions/choosing-ai-guardrails-legal","cat":"Blog"},{"title":"Best Prompt Observability Tools in 2026","desc":"The 6 best prompt observability tools in 2026, ranked on tracing, logging, latency and cost tracking, and debugging prompts running live in production.","href":"/blog/best-prompt-observability-tools-in-2026","cat":"Blog"},{"title":"How to Build a Coding Agent Harness From Scratch","desc":"What goes into a coding agent harness: the agent loop, tool registry, stop conditions, context management, and sandboxing, explained with real build examples.","href":"/blog/how-to-build-a-coding-agent-harness","cat":"Blog"},{"title":"Choosing AI Guardrails for Healthcare AI: 7 Questions to Ask","desc":"Seven questions for choosing an AI guardrails platform for healthcare AI: real-time PHI blocking, output-side scanning, hallucination checks, and self-hosting.","href":"/blog/choosing-x-tools-y-questions/choosing-ai-guardrails-healthcare","cat":"Blog"},{"title":"Best Prompt Registry Tools for AI Teams in 2026","desc":"A ranked guide to the 6 best prompt registry tools for AI teams in 2026, compared on central storage, versioning, labels, and fetching prompts at runtime.","href":"/blog/best-prompt-registry-tools-for-ai-teams-in-2026","cat":"Blog"},{"title":"Choosing the Right Content Moderation Tool for Generative AI Apps: 5 Questions to Ask","desc":"Five questions for choosing a content moderation tool for generative AI apps: output moderation, runtime blocking, custom policy, and multimodal coverage.","href":"/blog/choosing-x-tools-y-questions/choosing-content-moderation-tool","cat":"Blog"},{"title":"Choosing the Right Prompt Injection Defense Tool for AI Agents: 8 Factors to Compare","desc":"Eight factors for choosing a prompt injection defense tool for AI agents: catching indirect injection, blocking at runtime, constraining tools, securing MCP.","href":"/blog/choosing-x-tools-y-questions/choosing-prompt-injection-defense-tool","cat":"Blog"},{"title":"Best Prompt Deployment Platforms for Production LLM Apps in 2026","desc":"Compare the 6 best prompt deployment platforms for production LLM apps in 2026 on labels, rollback, runtime fetching, and safe, controlled prompt releases.","href":"/blog/best-prompt-deployment-platforms-for-production-llm-apps-in-2026","cat":"Blog"},{"title":"Choosing the Right PII Redaction Tool for LLM Applications: 7 Questions Before You Buy","desc":"Seven questions for choosing a PII redaction tool for LLM applications: runtime redaction, contextual detection, whole-path coverage, and output leaks.","href":"/blog/choosing-x-tools-y-questions/choosing-pii-redaction-tool","cat":"Blog"},{"title":"Choosing the Right Online Evaluation Platform for Live Production Traffic: 6 Critical Factors","desc":"Six factors for choosing an online evaluation platform for live production traffic: scoring on the trace, sampling at scale, quality alerts, and root cause.","href":"/blog/choosing-x-tools-y-questions/choosing-online-evaluation-platform","cat":"Blog"},{"title":"Best Prompt Versioning Tools for Production AI Teams in 2026","desc":"The 7 best prompt versioning tools for production AI teams in 2026, ranked on version history, labels, rollback, and testing prompts before you ship it.","href":"/blog/best-prompt-versioning-tools-for-production-ai-teams-in-2026","cat":"Blog"},{"title":"Choosing the Right Offline Evaluation Tool for Pre-Deployment Testing: 5 Questions to Ask","desc":"Five questions for choosing an offline evaluation tool for pre-deployment testing: run on a test set, build one without data, catch regressions, and gate CI.","href":"/blog/choosing-x-tools-y-questions/choosing-offline-evaluation-tool","cat":"Blog"},{"title":"Building Custom Eval Metrics for Industry-Specific Evaluation: 8 Questions to Ask","desc":"Eight questions for building custom eval metrics for domain-specific evaluation: author and calibrate your own metrics, then run them in CI and production.","href":"/blog/choosing-x-tools-y-questions/building-custom-eval-metrics","cat":"Blog"},{"title":"Best Prompt Chaining Tools for LLM Workflows in 2026","desc":"See the 6 best prompt chaining tools for LLM workflows in 2026, ranked on per-link versioning, tracing, evaluation, and debugging each step of the chain.","href":"/blog/best-prompt-chaining-tools-for-llm-workflows-in-2026","cat":"Blog"},{"title":"Choosing the Right LLM Evaluation SDK for Engineering Teams: 6 Critical Factors","desc":"Six factors for choosing an LLM evaluation SDK for engineering teams: evals in code and CI, no-ground-truth scoring, and local-plus-judge routing.","href":"/blog/choosing-x-tools-y-questions/choosing-llm-evaluation-sdk","cat":"Blog"},{"title":"Best Open-Source Prompt Management Tools in 2026","desc":"Compare the 6 best open-source prompt management tools of 2026 on self-hosting, licensing, versioning, and built-in evaluation you run on your own stack.","href":"/blog/best-open-source-prompt-management-tools-in-2026","cat":"Blog"},{"title":"Best Prompt Management Tools for Startups in 2026","desc":"The 6 best prompt management tools for startups in 2026, ranked on free tiers, quick setup, versioning, and built-in evaluation for small, fast teams.","href":"/blog/best-prompt-management-tools-for-startups-in-2026","cat":"Blog"},{"title":"Choosing the Right Hallucination Detection Tool for Production AI: 8 Factors to Compare","desc":"Eight factors for choosing a hallucination detection tool for production AI, from catching errors with no reference answer to blocking bad responses live instead of flagging them once they've already reached users.","href":"/blog/choosing-x-tools-y-questions/choosing-hallucination-detection-tool","cat":"Blog"},{"title":"Best Prompt Management Tools for CrewAI Agents in 2026","desc":"The 6 best prompt management tools for CrewAI agents in 2026, ranked on integration, per-agent versioning, tracing, and built-in evaluation of prompts.","href":"/blog/best-prompt-management-tools-for-crewai-agents-in-2026","cat":"Blog"},{"title":"Best Prompt Engineering Tools for Production LLM Apps in 2026","desc":"The 6 best prompt engineering tools for production LLM apps in 2026, ranked on versioning, evaluation, CI gates, tracing, and quick rollback in production.","href":"/blog/best-prompt-engineering-tools-for-production-llm-apps-in-2026","cat":"Blog"},{"title":"Best Prompt IDE Tools in 2026","desc":"The 6 best prompt IDE tools in 2026, ranked on the editor, variables, side-by-side testing, versioning, and evaluation together in a single workspace.","href":"/blog/best-prompt-ide-tools-in-2026","cat":"Blog"},{"title":"Best Prompt Registry Platforms for MLOps Teams in 2026","desc":"Compare the 6 best prompt registry platforms for MLOps teams in 2026 on prompt versioning, model registry fit, CI/CD, and production deployment workflows.","href":"/blog/best-prompt-registry-platforms-for-mlops-teams-in-2026","cat":"Blog"},{"title":"Best Git-Based Prompt Management Platforms in 2026","desc":"Compare the 6 best git-based prompt management platforms of 2026 on version control, branching, diffs, and code-first prompt workflows for engineering teams.","href":"/blog/best-git-based-prompt-management-platforms-in-2026","cat":"Blog"},{"title":"Best Collaborative Prompt Management Platforms for Product Teams in 2026","desc":"The 6 best collaborative prompt management platforms for product teams in 2026, ranked on shared editing, roles, versioning, and safe, reviewed releases.","href":"/blog/best-collaborative-prompt-management-platforms-for-product-teams-in-2026","cat":"Blog"},{"title":"Best Prompt Iteration Tools in 2026","desc":"The 7 best prompt iteration tools in 2026, ranked on fast editing, versioning, evaluation, and comparing prompt changes before you ship them to users.","href":"/blog/best-prompt-iteration-tools-in-2026","cat":"Blog"},{"title":"Choosing the Right RAG Observability Tool for Pipelines: 7 Critical Factors","desc":"A 7-factor guide to choosing a RAG observability tool: retriever-span tracing, retrieval-vs-generation attribution, and ground-truth-free scoring.","href":"/blog/choosing-x-tools-y-questions/choosing-rag-observability-tool","cat":"Blog"},{"title":"Best Prompt Experimentation Platforms in 2026","desc":"Compare the 6 best prompt experimentation platforms of 2026 on A/B testing, versioning, evaluation, and comparing prompt variants side by side at scale.","href":"/blog/best-prompt-experimentation-platforms-in-2026","cat":"Blog"},{"title":"Best Prompt Management Tools for LlamaIndex Apps in 2026","desc":"See the 6 best prompt management tools for LlamaIndex apps in 2026, ranked on integration, versioning, tracing, and evaluation for your RAG app prompts.","href":"/blog/best-prompt-management-tools-for-llamaindex-apps-in-2026","cat":"Blog"},{"title":"Best Tools for Creating System Prompts in 2026","desc":"The 6 best tools for creating system prompts in 2026, ranked on structured authoring, versioning, testing, and maintaining a strong system prompt over time.","href":"/blog/best-tools-for-creating-system-prompts-in-2026","cat":"Blog"},{"title":"Choosing the Right Eval-Driven Development Platform for AI Engineering Teams: 8 Factors to Compare","desc":"Eight factors for choosing an eval-driven development platform for AI teams: evals as spec, CI gating, dev-prod parity, regression loops, and optimization.","href":"/blog/choosing-x-tools-y-questions/choosing-eval-driven-development-platform","cat":"Blog"},{"title":"Best Prompt Management Tools for LangChain Apps in 2026","desc":"See the 6 best prompt management tools for LangChain apps in 2026, ranked on native integration, versioning, tracing, and evaluation across your chains.","href":"/blog/best-prompt-management-tools-for-langchain-apps-in-2026","cat":"Blog"},{"title":"Best Prompt Governance Platforms for Enterprise AI in 2026","desc":"The 6 best prompt governance platforms for enterprise AI in 2026, ranked on access control, audit logs, approval workflows, and clear policy enforcement.","href":"/blog/best-prompt-governance-platforms-for-enterprise-ai-in-2026","cat":"Blog"},{"title":"Best Prompt Management Tools with Built-In Evaluation in 2026","desc":"The 6 best prompt management tools with built-in evaluation in 2026, ranked on evaluator libraries, version-tied scoring, regression testing, and CI checks before release.","href":"/blog/best-prompt-management-tools-with-built-in-evaluation-in-2026","cat":"Blog"},{"title":"Best Prompt Orchestration Platforms for AI Agents in 2026","desc":"Compare the 6 best prompt orchestration platforms for AI agents in 2026 on per-step versioning, tracing, evaluation, and automated prompt optimization.","href":"/blog/best-prompt-orchestration-platforms-for-ai-agents-in-2026","cat":"Blog"},{"title":"Best Prompt Management Tools for AutoGen Agents in 2026","desc":"See the 6 best prompt management tools for AutoGen agents in 2026, ranked on per-agent versioning, runtime fetching, tracing, and built-in evaluation.","href":"/blog/best-prompt-management-tools-for-autogen-agents-in-2026","cat":"Blog"},{"title":"Best Prompt Management Tools for LangGraph Apps in 2026","desc":"The 6 best prompt management tools for LangGraph apps in 2026, ranked on per-node versioning, tracing, evaluation, and safe deployment of node prompts.","href":"/blog/best-prompt-management-tools-for-langgraph-apps-in-2026","cat":"Blog"},{"title":"Choosing the Right AI Gateway for Built-In Guardrails: 8 Factors to Compare","desc":"Eight factors for choosing an AI gateway with built-in guardrails: inline enforcement, scanner coverage, output and tool-call guarding, and self-hosting.","href":"/blog/choosing-x-tools-y-questions/choosing-ai-gateway-guardrails","cat":"Blog"},{"title":"Error Localization for Debugging LLM Apps: 7 Capabilities to Look For in an Eval Platform","desc":"Seven capabilities to look for in error localization when choosing an LLM eval platform: field-level attribution, span-tied failures, reasoning, and production fit.","href":"/blog/choosing-x-tools-y-questions/choosing-error-localization-tool","cat":"Blog"},{"title":"Choosing the Right Ground-Truth-Free Evaluation Tool for Teams Without Labeled Data: 5 Questions to Ask","desc":"Five questions for choosing a ground-truth-free eval tool with no labeled data: reference-free scoring, judge calibration, error localization, production fit.","href":"/blog/choosing-x-tools-y-questions/choosing-ground-truth-free-eval-tool","cat":"Blog"},{"title":"Choosing the Right Prompt Optimization Tool for Production Teams: 8 Questions to Ask","desc":"Eight questions for choosing a prompt optimization tool for production teams: automated search over your own eval scores, named algorithms, open source, convergence, and CI/CD fit.","href":"/blog/choosing-x-tools-y-questions/choosing-prompt-optimization-tool","cat":"Blog"},{"title":"Agent Runtime Guardrails in 2026: The Tool-Call Scanners Most Stacks Skip","desc":"PII and toxicity scanners never see the tool call. Agent runtime guardrails (tool permissions, MCP security, system-prompt protection) catch what they miss.","href":"/blog/agent-runtime-guardrails","cat":"Blog"},{"title":"Automatic Prompt Optimization in 2026: How Textual Gradients, Genetic Search, and Meta-Prompts Actually Work","desc":"Automatic prompt optimization explained: textual gradients (ProTeGi), score trajectories (OPRO), genetic evolution (GEPA), meta-prompting, and how to pick one.","href":"/blog/automatic-prompt-optimization","cat":"Blog"},{"title":"DSPy Optimizers Explained in 2026: BootstrapFewShot, MIPROv2, COPRO, and GEPA","desc":"A practitioner's comparison of DSPy optimizers: how BootstrapFewShot, MIPROv2, COPRO, and GEPA differ, and a ladder for picking the right one.","href":"/blog/dspy-optimizers-explained","cat":"Blog"},{"title":"Falcon AI in 2026: The Platform-Native Copilot That Operates Your Eval Stack","desc":"A generic chatbot answers questions about your data. Falcon AI runs the eval, drills the trace, and files the ticket, with 300+ tools and page context.","href":"/blog/falcon-ai-copilot","cat":"Blog"},{"title":"Your LLM Eval Failed. Which Input Broke It? Field-Level Eval Attribution in 2026","desc":"A pass/fail eval score says something broke, not what. Field-level eval attribution pins the failure to the exact input: context, question, or output.","href":"/blog/llm-eval-error-localization","cat":"Blog"},{"title":"Multimodal LLM-as-a-Judge in 2026: How to Evaluate Images and Audio Without Ground Truth","desc":"Text-only evals never check the image. How a multimodal LLM-as-a-judge scores image-text alignment, generated images, and audio, with no reference.","href":"/blog/multimodal-llm-as-a-judge","cat":"Blog"},{"title":"Inside Observe: The Six Surfaces of Production Agent Observability in 2026","desc":"Production observability has to answer six questions. Here is the Observe surface for each: sessions, users, trace evals, dashboards, alerts, and voice.","href":"/blog/observe-surfaces-tour","cat":"Blog"},{"title":"Production Replay Testing in 2026: How to Simulate Real Sessions, Traces, and Calls","desc":"Synthetic test cases can't reproduce the bug a real user hit. Production replay reruns the exact session, trace, or voice call against your fixed agent.","href":"/blog/production-replay-testing","cat":"Blog"},{"title":"Scenarios vs Synthetic Data in 2026: Why Testing an Agent Isn't Generating Rows","desc":"Synthetic data is static rows you score once. A scenario is a multi-turn conversation your agent has to navigate. Here is the difference and when each fits.","href":"/blog/scenario-vs-synthetic-data","cat":"Blog"},{"title":"Trace-Native Evaluation in 2026: Score the Whole Trace, Skip the Data Mapping","desc":"Most eval loops export logs, build a dataset, map columns. Trace-native evaluation attaches the score to the span itself and runs on production traces.","href":"/blog/trace-native-evaluation","cat":"Blog"},{"title":"Choosing the Right Compliance Monitoring Tool for AI Governance: 7 Questions to Ask","desc":"Seven questions for choosing an AI governance compliance monitoring tool: runtime guardrails, PII and prompt-injection enforcement, audit trails, self-hosting.","href":"/blog/choosing-x-tools-y-questions/choosing-compliance-monitoring-tool","cat":"Blog"},{"title":"How we redesigned futureagi.com: a starship for AI in production","desc":"The customer email that started it, the starship metaphor we almost cut, the hyperspace footer, the handbook fight, and the surfaces we honestly haven't finished yet.","href":"/blog/redesigning-futureagi-2026","cat":"Blog"},{"title":"Choosing the Right Voice Agent Testing Platform: 14 Questions to Ask","desc":"Fourteen questions for choosing a voice agent testing platform: persona-driven simulation, multi-turn coverage, adversarial users, tone and audio scoring, and turning failures into fixes.","href":"/blog/choosing-x-tools-y-questions/choosing-voice-agent-testing-platform","cat":"Blog"},{"title":"Choosing the Right LLM Evaluation Platform for Production: 10 Critical Factors","desc":"A 10-factor checklist for choosing an LLM evaluation platform for production: ground-truth-free scoring, error localization, online eval, data handling, and more.","href":"/blog/choosing-x-tools-y-questions/choosing-llm-evaluation-platform-production","cat":"Blog"},{"title":"Best 5 Literal AI Alternatives in 2026 (Migration Guide)","desc":"Literal AI's hosted platform was discontinued. This migration guide ranks five alternatives and shows how to move traces, datasets, and prompts off it.","href":"/blog/best-literal-ai-alternatives-2026","cat":"Blog"},{"title":"Best 5 Parea AI Alternatives in 2026","desc":"Five Parea AI alternatives scored on eval-catalog depth, logs-capped pricing, optimizer loops, guardrails, and team scale, and what each fixes.","href":"/blog/best-parea-ai-alternatives-2026","cat":"Blog"},{"title":"Best 5 RagaAI Alternatives in 2026","desc":"Five RagaAI alternatives scored on eval-judge depth, optimizer loops, gateway and guardrails, self-host ops burden, vendor maturity, and what each fixes.","href":"/blog/best-ragaai-alternatives-2026","cat":"Blog"},{"title":"Evaluating Pydantic AI Agents That Use MCP Tools (2026)","desc":"Evaluate Pydantic AI agents that call MCP tools in 2026: per-typed-output rubrics, tool-call argument fidelity, MCP security checks, dependency invariants.","href":"/blog/evaluating-pydantic-ai-mcp-agents-2026","cat":"Blog"},{"title":"Evaluating Vector Database Recall Quality in 2026","desc":"Vendor vector-DB benchmarks are theater. ANN-vs-exact-knn recall on your vectors plus p99 under your filter cardinality is the eval that decides prod.","href":"/blog/evaluating-vector-database-recall-quality-2026","cat":"Blog"},{"title":"Future AGI vs Parea AI 2026: Closed Loop vs Annotation-First","desc":"Future AGI vs Parea AI scored on tracing, evaluation, prompt management, simulation, security, and DX. Honest verdict and May 2026 pricing.","href":"/blog/future-agi-vs-parea-ai-2026","cat":"Blog"},{"title":"Future AGI vs RagaAI 2026: Closed Loop vs Multi-Domain","desc":"Future AGI vs RagaAI scored on tracing, evaluation, prompt management, simulation, security, and DX. Honest verdict and May 2026 pricing.","href":"/blog/future-agi-vs-ragaai-2026","cat":"Blog"},{"title":"LLM Eval with Shadow Traffic and Canary Deployment in 2026","desc":"Shadow is not canary. Mirror routing with no user effect vs percentage routing with rollback. Score-attached traffic, ACC patterns, gotchas.","href":"/blog/llm-eval-shadow-traffic-canary-2026","cat":"Blog"},{"title":"How to Evaluate RAG Applications in CI/CD Pipelines (2026)","desc":"RAG eval in CI/CD without theatre: the cheap-fast-significant triangle, statistical gating, sharded parallelism, classifier cascades, production bridge.","href":"/blog/evaluate-rag-applications-cicd-2026","cat":"Blog"},{"title":"Evaluating Azure OpenAI LLM Apps in 2026","desc":"Azure OpenAI eval has three Azure-specific axes: deployment-name drift, region-pinning, and Content Safety precision on benign queries. Here's the pattern.","href":"/blog/evaluating-azure-openai-llm-apps-2026","cat":"Blog"},{"title":"How to Jailbreak LLMs (Defender's Guide): A Step-by-Step Walkthrough","desc":"Defender's walkthrough of LLM jailbreak techniques in 2026: role-play, encoding, multi-turn drift, indirect injection. Each attack mapped to a guardrail.","href":"/blog/llm-jailbreak-step-by-step-2026","cat":"Blog"},{"title":"Evaluating AWS Bedrock Agents in 2026","desc":"Bedrock's built-in eval is dev-loop only. Score action-group correctness, KB retrieval quality, and guardrail precision/recall on every release.","href":"/blog/evaluating-aws-bedrock-agents-2026","cat":"Blog"},{"title":"Gemini 3.5 Flash: The Numbers Behind the May 2026 Launch","desc":"Gemini 3.5 Flash dropped today at Google I/O 2026. The 8 benchmark numbers that matter, $1.50/$9 pricing breakdown, and what to instrument before you swap.","href":"/blog/gemini-3-5-flash-launch-2026","cat":"Blog"},{"title":"LLM Eval Budget Allocation and Prioritization in 2026","desc":"Eval budget is four knobs: rubric coverage, dataset size, judge tier, refresh cadence. Priority order that maximizes signal per dollar, with a 90-day plan.","href":"/blog/llm-eval-budget-allocation-prioritization-2026","cat":"Blog"},{"title":"Evaluating LLM Personas and Style Drift (2026)","desc":"Persona eval is two problems: per-turn tone and cross-turn drift. The 2026 playbook for tone rubrics plus a persona-stability score across turns.","href":"/blog/evaluating-llm-personas-style-2026","cat":"Blog"},{"title":"Evaluating LLM Translation Quality (2026)","desc":"BLEU is dead for LLM translation. The 2026 stack: COMET + LLM-as-judge fluency/adequacy rubrics + per-language-pair calibration. With code and thresholds.","href":"/blog/evaluating-llm-translation-quality-2026","cat":"Blog"},{"title":"LLM Pricing and Cost Comparison Guide for 2026","desc":"Sticker per-token price is the wrong unit. Use effective cost = sticker x (1 - cache_hit) x prompt_ratio. 2026 methodology, five lanes, per-provider knobs.","href":"/blog/llm-pricing-cost-comparison-2026","cat":"Blog"},{"title":"Best 5 Pydantic AI Alternatives in 2026","desc":"Five Pydantic AI alternatives on multi-agent depth, language reach, observability without Logfire, optimizer. What each actually fixes past type-system.","href":"/blog/best-pydantic-ai-alternatives-2026","cat":"Blog"},{"title":"Evaluating LiteLLM Multi-Provider Apps in 2026","desc":"How to evaluate LiteLLM-routed apps: paired comparison across providers on your data, tool-call parity, latency parity, and the gateway alternative.","href":"/blog/evaluating-litellm-multi-provider-2026","cat":"Blog"},{"title":"15 Common LLM Evaluation Mistakes Teams Make in 2026","desc":"The 15 LLM evaluation mistakes the Future AGI team sees in customer engagements, each with a vignette and the concrete primitive that prevents it.","href":"/blog/llm-eval-common-mistakes-2026","cat":"Blog"},{"title":"The Alignment Paradox: A 2026 Practitioner Reading","desc":"Helpful and harmless trade off. Labs that pretend otherwise are training to a benchmark, not behavior. The alignment paradox in mid-2026.","href":"/blog/alignment-paradox-llm-safety-2026","cat":"Blog"},{"title":"Best 5 AI Gateways to Cache Claude Code Calls in 2026","desc":"Five AI gateways scored on caching Claude Code calls in 2026: cross-developer cache scope, semantic-match thresholds, hit-rate, TTL, what each misses.","href":"/blog/best-ai-gateways-claude-code-caching-2026","cat":"Blog"},{"title":"LLM Eval Golden Set Design: A 2026 Engineering Guide","desc":"Build a four-bucket golden set (production sample, adversarial, edge cases, failure replays) so a CI eval gate actually proves something about production.","href":"/blog/llm-eval-golden-set-design-2026","cat":"Blog"},{"title":"AI Gateway for Codex CLI in 2026: The Playbook","desc":"Wrap OpenAI Codex CLI in an AI gateway for per-developer budgets, per-call audit trail, and provider flexibility, without changing the CLI command.","href":"/blog/ai-gateway-codex-cli-governance-cost-provider-flexibility-2026","cat":"Blog"},{"title":"Future AGI vs LiteLLM in 2026: Self-Improving Runtime vs OSS Python Proxy","desc":"Future AGI vs LiteLLM scored on routing, observability, cost attribution, security, deployment, DX. Honest verdict, March 2026 PyPI compromise context.","href":"/blog/future-agi-vs-litellm-2026","cat":"Blog"},{"title":"Future AGI vs Portkey in 2026: Self-Improving Runtime vs Hosted Gateway","desc":"Future AGI vs Portkey scored on routing, observability, cost attribution, security, deployment, DX. Why FAGI wins the self-improving loop, post-PANW note.","href":"/blog/future-agi-vs-portkey-2026","cat":"Blog"},{"title":"Running Claude Code with OpenAI Models in 2026: A Gateway Setup Guide","desc":"Run Claude Code against OpenAI GPT-5 and GPT-4 via a translation gateway in 2026: setup, ENV vars, config, then five gateways scored.","href":"/blog/running-claude-code-with-openai-models-2026","cat":"Blog"},{"title":"Custom Voice Evaluator Authoring in 2026: The In-Product Agent Workflow","desc":"Author custom voice evaluators in 2026 two ways: in-product agent that proposes rubrics from traces, plus code that extends Evaluator class.","href":"/blog/custom-voice-evaluator-authoring-2026","cat":"Blog"},{"title":"Evaluating RAG Chunking Strategies in 2026","desc":"Chunking is a domain question. Fixed-size loses on legal, semantic wins prose, clause-level wins contracts, late-interaction wins code.","href":"/blog/evaluating-rag-chunking-strategies-2026","cat":"Blog"},{"title":"Top 5 Tools for Claude Code Cost Management in 2026","desc":"Five tools for Claude Code cost management in 2026: four gateways, the native Anthropic dashboard, and a FinOps platform, scored on chargeback, caps.","href":"/blog/top-5-tools-claude-code-cost-management-2026","cat":"Blog"},{"title":"Best AI Gateway for Lovable AI App Generator in 2026","desc":"Five AI gateways scored on Lovable AI App Generator in 2026: B2B2C per-customer attribution, per-project caps, iteration traces, brand-safety, gaps.","href":"/blog/best-ai-gateway-lovable-ai-app-generator-2026","cat":"Blog"},{"title":"Breaking Gemini (Defender's View): Runtime Defense for Google's Models","desc":"Gemini wins on single-turn refusal precision, loses on multi-turn Crescendo and context drift. Defender's read on 2.5 and 3, the layer builders owe.","href":"/blog/breaking-gemini-defender-analysis-2026","cat":"Blog"},{"title":"Intent Classification Evaluation Pipeline (2026)","desc":"Per-intent precision-recall, escalation accuracy, OOD detection, and drift gates for the LLM router that decides which pipeline runs.","href":"/blog/intent-classification-evaluation-pipeline-2026","cat":"Blog"},{"title":"Best 5 AI Gateways to Manage Cursor Spend Across Teams in 2026","desc":"Five AI gateways scored on Cursor team spend in 2026: per-dev chargeback, per-repo budgets, SSO attribution, BYOK virtual keys, where each falls short.","href":"/blog/best-ai-gateways-cursor-spend-teams-2026","cat":"Blog"},{"title":"Best 5 Voice AI Simulation Tools for CX AI Applications in 2026","desc":"Five voice AI simulation tools compared for CX, IVR upgrades, outbound TCPA, multi-turn refunds. FCC AI-voice, recording consent, FCRA Reg F.","href":"/blog/best-cx-voice-ai-simulation","cat":"Blog"},{"title":"Best Cybersecurity AI Evaluation Platforms in 2026","desc":"Cybersecurity AI eval in 2026: five platforms scored on red-team rubric, false-positive floor, prompt-injection scanner integration. FAGI, Galileo, Lakera.","href":"/blog/best-cybersecurity-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best Education AI Evaluation Platforms in 2026","desc":"Education AI eval in 2026: five platforms scored on COPPA + FERPA + pedagogical-correctness rubrics. FAGI, Galileo Luna-2, Braintrust, Khanmigo, on-prem.","href":"/blog/best-education-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best 5 AI Evaluation Platforms for Government in 2026: FedRAMP, IL5, NIST AI RMF, Section 508","desc":"Five AI evaluation platforms scored for public-sector AI on FedRAMP, IL5, StateRAMP, NIST AI RMF, air-gap, and Section 508. May 2026.","href":"/blog/best-government-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for Healthcare AI Applications in 2026","desc":"Five AI guardrails compared for healthcare: clinical decision, ambient scribes, prior auth, portal chatbots. HIPAA, FDA SaMD, EU AI Act 14.","href":"/blog/best-healthcare-ai-guardrails-2026","cat":"Blog"},{"title":"Best HR AI Evaluation Platforms in 2026","desc":"HR AI eval in 2026: five platforms scored on demographic bias detection, per-decision audit, impact-ratio reporting. FAGI, Galileo, Braintrust, Holistic.","href":"/blog/best-hr-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best 5 AI Evaluation Tools for Manufacturing AI Applications in 2026","desc":"Five AI eval platforms for manufacturing, predictive maintenance, defect, MES copilots, safety docs. ISO 9001, OSHA 5(a)(1), EU 2023/1230, CMMC, NIST AI.","href":"/blog/best-manufacturing-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for Retail AI Applications in 2026","desc":"Five AI guardrails platforms for retail: returns chatbots, recommendation engines, PDP generation, dynamic pricing, conversational commerce. FTC, PCI-DSS.","href":"/blog/best-retail-ai-guardrails-2026","cat":"Blog"},{"title":"Future AGI vs LangSmith in 2026: Self-Improving Runtime vs Hosted Observability","desc":"Future AGI vs LangSmith on tracing, evaluation, prompt management, deployment, security, DX. Honest verdict, May 2026, why only one closes the loop.","href":"/blog/future-agi-vs-langsmith-2026","cat":"Blog"},{"title":"Best 7 AI Gateways for Agentic AI in 2026","desc":"Seven AI gateways for agentic AI in 2026 scored on MCP plus A2A depth, 18+ built-in guardrails, OTel-native cost telemetry, 2026 trust cohort with sources.","href":"/blog/best-ai-gateways-agentic-ai","cat":"Blog"},{"title":"Best 5 AI Gateways for Cybersecurity in 2026: Prompt Injection Defense, Tenant Isolation, and SOC 2","desc":"Five AI gateways for cybersecurity in 2026 scored on prompt injection defense, tenant isolation, OWASP LLM Top 10, MITRE ATLAS, SOC 2.","href":"/blog/best-ai-gateways-cybersecurity-2026","cat":"Blog"},{"title":"Best CX AI Observability Platforms in 2026: 5 Picks","desc":"Five CX AI observability platforms scored on conversation-trace inspection, escalation-event capture, and CSAT/NPS join to Zendesk and Intercom ticket IDs.","href":"/blog/best-cx-ai-observability-2026","cat":"Blog"},{"title":"Best 5 RAG Evaluation Tools for Customer Support AI Applications in 2026","desc":"Five RAG eval tools for customer support, copilot, KB chatbot, billing agent. FTC Op AI Comply, Moffatt v. Air Canada, EU AI Act Art 50.","href":"/blog/best-cx-rag-evaluation-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for Cybersecurity AI Applications in 2026","desc":"Five AI guardrails compared for cybersecurity: SOC copilots, threat-intel RAG, SIEM LLMs, IR chatbots. NIST AI RMF, OWASP, MITRE ATLAS.","href":"/blog/best-cybersecurity-ai-guardrails-2026","cat":"Blog"},{"title":"Best 5 AI Observability Tools for Fintech AI Applications in 2026","desc":"Five fintech AI observability platforms scored on per-decision spans, immutable audit, SOC 2 + PCI-DSS, FFIEC / SR 11-7 model risk, EU DORA alignment.","href":"/blog/best-fintech-ai-observability-2026","cat":"Blog"},{"title":"Best 5 RAG Evaluation Tools for Fintech AI Applications in 2026","desc":"Five RAG eval tools for fintech: advisor copilots, KYC RAG, credit-decisioning RAG, regulatory research. NYDFS, FINRA, SEC 17a-4, CFPB audit covered.","href":"/blog/best-fintech-rag-evaluation-2026","cat":"Blog"},{"title":"Best 5 AI Observability Tools for Healthcare AI Applications in 2026","desc":"Five healthcare AI observability platforms scored on HIPAA trace ingestion, §164.312(b) retention, per-clinician access, BAA-boundary integrity. May 2026.","href":"/blog/best-healthcare-ai-observability-2026","cat":"Blog"},{"title":"Best 5 RAG Evaluation Tools for Healthcare AI Applications in 2026","desc":"Five RAG evaluation tools for healthcare: clinical decision support, ambient scribes, prior auth, medical coding. HIPAA, FDA SaMD, Cures Act, EU AI Act.","href":"/blog/best-healthcare-rag-evaluation-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for HR AI Applications in 2026","desc":"Five AI guardrails for HR: resume screening, AI interviews, internal mobility, recruiter copilots. NYC AEDT, EEOC 4/5 rule, FCRA, ADEA, EU AI Act III.","href":"/blog/best-hr-ai-guardrails-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for Insurance AI Applications in 2026","desc":"Five AI guardrails for insurance: underwriting, claims triage, fraud, copilots, CS chatbots, renewal pricing. NAIC, CO SB 21-169, NY DFS CL 7, ACA §1557.","href":"/blog/best-insurance-ai-guardrails-2026","cat":"Blog"},{"title":"Best 5 AI Observability Tools for Insurance AI Applications in 2026","desc":"Five AI observability tools for insurance: underwriting copilots, claims triage, fraud detection. NAIC, Colorado SB 21-169, NY DFS CL 7, GLBA, ACA §1557.","href":"/blog/best-insurance-ai-observability-2026","cat":"Blog"},{"title":"Best 5 RAG Evaluation Tools for Insurance AI Applications in 2026","desc":"Five RAG evaluation tools for insurance: underwriting, claims triage, fraud detection, agent copilots. NAIC, Colorado SB 21-169, NY DFS CL 7, NY Reg 187.","href":"/blog/best-insurance-rag-evaluation-2026","cat":"Blog"},{"title":"Best 5 AI Observability Tools for Legal AI Applications in 2026","desc":"Five AI observability tools compared for legal research, contract review, e-discovery. ABA Rules 1.1/1.6, Mata v. Avianca, FRCP 26(g).","href":"/blog/best-legal-ai-observability-2026","cat":"Blog"},{"title":"Best 5 RAG Evaluation Tools for Legal AI Applications in 2026","desc":"Five RAG evaluation tools for legal: brief drafting, contract review, e-discovery. ABA Model Rules 1.1/1.6/3.3/5.3, Mata v. Avianca, FRCP 11/26(g).","href":"/blog/best-legal-rag-evaluation-2026","cat":"Blog"},{"title":"Best 5 AI Observability Tools for SaaS AI Applications in 2026","desc":"Five AI observability tools for SaaS: multi-tenant span isolation, per-tenant cost attribution, LangChain fan-out. SOC 2, GDPR, EU AI Act. May 2026.","href":"/blog/best-saas-ai-observability-tools","cat":"Blog"},{"title":"Best 5 AI Observability Tools for Retail AI Applications in 2026","desc":"Five AI observability tools for retail, rec engines, PDP gen, merchandising copilots, CS chatbots, AI reviews. FTC §5, Op AI Comply, PCI-DSS, ADA.","href":"/blog/best-retail-ai-observability-2026","cat":"Blog"},{"title":"Best 5 RAG Evaluation Tools for SaaS AI Applications in 2026","desc":"Five RAG eval tools for SaaS, multi-tenant copilot, KB chatbot, search/Q&A, agentic IDE. SOC 2, GDPR Art 28, EU AI Act Art 50, FTC Op AI Comply.","href":"/blog/best-saas-rag-evaluation-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for LLM Cost Optimization in 2026","desc":"Five AI gateways for LLM cost optimization in 2026 scored on the five-layer cost stack, 18+ guardrails, OTel-native cost telemetry, and 2026 trust cohort.","href":"/blog/best-ai-gateways-cost-optimization","cat":"Blog"},{"title":"Evaluating MCP Servers for Security (2026)","desc":"Evaluate MCP servers for security in 2026: tool-description injection, tool-result tampering, sandbox escape, cross-tenant isolation. Four eval checks.","href":"/blog/evaluating-mcp-servers-security-2026","cat":"Blog"},{"title":"LLM Eval for Startups in 2026: A Lean Quality Discipline","desc":"How an 8-engineer startup ships production LLM eval without a dedicated team: seven principles, five-engineer rollout, the FAGI primitives that scale.","href":"/blog/llm-eval-for-startups-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for Education AI Applications in 2026","desc":"Five AI guardrails platforms for education: K-12 tutoring, curriculum copilots, grading, student-records agents, IEP copilots. FERPA, COPPA, PPRA, CIPA.","href":"/blog/best-education-ai-guardrails-2026","cat":"Blog"},{"title":"The Definitive Guide to Synthetic Data Generation with LLMs (2026)","desc":"The 2026 reference: three generation patterns (persona, taxonomy-stratified, evolution), the filter that survives, calibration against real, use cases.","href":"/blog/definitive-guide-synthetic-data-generation-2026","cat":"Blog"},{"title":"LLM Eval vs Agent Protocol Evolution: Going Protocol-Neutral in 2026","desc":"MCP, A2A, OpenAI Responses, Realtime, Anthropic Tool Use, Google ADK. Six protocols changed agent eval. Keep your eval stack neutral.","href":"/blog/llm-eval-vs-mcp-protocol-evolution-2026","cat":"Blog"},{"title":"Automated Optimization for Agents in 2026: 5 Axes, Not 1","desc":"Automated optimization for agents in 2026 is five axes: system prompt, tool descriptions, retrieval config, few-shot bundle, model. Pick the right one.","href":"/blog/automated-optimization-for-agent-2026","cat":"Blog"},{"title":"Best 5 AI Gateways to Monitor Claude Code Token Usage in 2026","desc":"Five AI gateways scored on Claude Code token monitoring in 2026: per-dev attribution, per-repo budgets, session traces, alerts, where each falls short.","href":"/blog/best-ai-gateways-claude-code-token-usage-2026","cat":"Blog"},{"title":"Best 5 Eyer AI Alternatives in 2026","desc":"Five Eyer AI alternatives on multi-language SDK coverage, self-host, gateway, optimizer reach. What each actually fixes outgrowing AI-monitoring-only.","href":"/blog/best-eyer-ai-alternatives-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for Fintech AI Applications in 2026","desc":"Five AI guardrails for fintech: fraud detection, credit, KYC, trading. NYDFS Part 500 §500.13, FINRA Rule 3110, SEC 15c3-5, EU AI Act Art 14, DORA.","href":"/blog/best-fintech-ai-guardrails-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for Legal AI Applications in 2026","desc":"Five AI guardrails for legal: brief drafting, contract review, legal research, e-discovery. ABA Model Rules 1.1/1.6/3.3/5.3, Mata v Avianca, FRCP 11.","href":"/blog/best-legal-ai-guardrails-2026","cat":"Blog"},{"title":"Weights & Biases Alternatives in 2026: 7 Platforms Compared","desc":"FutureAGI closes the self-improving loop; MLflow, Comet, Neptune, Langfuse, Braintrust, ClearML ship the parts. 2026 W&B alternatives.","href":"/blog/best-weights-and-biases-alternatives-2026","cat":"Blog"},{"title":"Introducing ai-evaluation: Future AGI's Open-Source LLM Eval Library","desc":"Introducing ai-evaluation, Future AGI's Apache 2.0 Python and TypeScript library for LLM evaluation. 50+ metrics, AutoEval, streaming, multimodal.","href":"/blog/ai-evaluation-open-source-llm-evaluation-library","cat":"Blog"},{"title":"Best 5 AI Gateways for Embedding API Routing in 2026","desc":"Five AI gateways for embedding API routing 2026: provider breadth, dimension consistency, batch APIs, input-hash cache, model migration.","href":"/blog/best-ai-gateways-embedding-api-routing-2026","cat":"Blog"},{"title":"Best Fintech AI Evaluation Platforms in 2026","desc":"Fintech AI eval in 2026: five platforms scored on SOC 2 + PCI-DSS, financial-regulation rubrics, SR 11-7. FAGI, Galileo Luna-2, Braintrust, Datadog.","href":"/blog/best-fintech-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best 5 Fireworks AI Alternatives in 2026","desc":"Five Fireworks AI alternatives on inference performance, catalog depth, fine-tuning ergonomics. What each actually fixes for production LLM workloads.","href":"/blog/best-fireworks-ai-alternatives-2026","cat":"Blog"},{"title":"Best Healthcare AI Evaluation Platforms in 2026","desc":"Healthcare AI eval in 2026: five platforms scored on HIPAA + BAA, clinical-grade rubrics, audit-trace retention. FAGI, Galileo Luna-2, Braintrust, Datadog.","href":"/blog/best-healthcare-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best Insurance AI Evaluation Platforms in 2026","desc":"Insurance AI eval in 2026: five platforms scored on bias detection, factuality, and per-decision audit. FAGI, Galileo Luna-2, Braintrust, Datadog, on-prem.","href":"/blog/best-insurance-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best Legal AI Evaluation Platforms 2026: 5 Compared","desc":"Five legal AI eval platforms scored on four tests: clause-citation validity, privilege, jurisdiction-correct authority, refusal calibration. May 2026.","href":"/blog/best-legal-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Best 5 AI Evaluation Platforms for Retail AI Applications in 2026","desc":"Five AI eval platforms for retail, recs, search, CS chatbots, PDP gen, dynamic pricing, conversational commerce. FTC, Moffatt, PCI-DSS v4, GDPR Art 22.","href":"/blog/best-retail-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Red-Teaming Conversational AI: What Your Voice Agent Should Never Say in 2026","desc":"Red-team voice agents against 8 attack archetypes in 2026 with Future AGI Protect, ProtectFlash, named eval rubrics, and 1,200-call pre-launch coverage.","href":"/blog/red-teaming-conversational-ai-voice-agents-2026","cat":"Blog"},{"title":"Anatomy of a Voice Agent Analytics Dashboard in 2026","desc":"Walkthrough of a voice agent analytics dashboard: per-call drawer with 5 panels, SLO grid with 3 tiers, span/eval/tag flow, production-to-sim closed loop.","href":"/blog/voice-agent-analytics-dashboard-anatomy-2026","cat":"Blog"},{"title":"Voice Agent Regression Testing in CI/CD: A 2026 Engineering Guide","desc":"Wire voice agent regression tests into GitHub Actions and GitLab CI: golden conversations, three-layer testing, deploy gates, FAGI evals.","href":"/blog/voice-agent-regression-testing-ci-cd-2026","cat":"Blog"},{"title":"Voice Agent Simulation: A 2026 Engineering Guide","desc":"Engineer voice agent simulation: 18 personas, auto-generated branching scenarios, four-step test wizard, Error Localization, programmatic eval API.","href":"/blog/voice-agent-simulation-2026-guide","cat":"Blog"},{"title":"Best LLM Cost Tracking Tools in 2026: 8 Compared","desc":"Future AGI, Helicone, Langfuse, OpenRouter, Portkey, LangSmith, Datadog, and CloudZero compared on per-trace, per-developer LLM cost attribution.","href":"/blog/best-llm-cost-tracking-tools-2026","cat":"Blog"},{"title":"Best LLMs of May 2026: Top Closed-Source, Open-Weight, Multimodal, and Coding Picks","desc":"Best LLMs May 2026: compare GPT-5.5, Claude Opus 4.7, Gemini 3.1 Pro, and DeepSeek V4 across coding, agents, multimodal, cost, and open weights.","href":"/blog/best-llms-may-2026","cat":"Blog"},{"title":"Best Voice AI Models in May 2026: STT, TTS, and Voice Agent Stack","desc":"Best Voice AI May 2026: compare Deepgram, Cartesia, ElevenLabs, Retell, and Vapi for STT, TTS, latency budgets, and production voice agents.","href":"/blog/best-voice-ai-may-2026","cat":"Blog"},{"title":"Future AGI vs OpenRouter in 2026: Self-Improving Runtime vs Hosted Router","desc":"Future AGI vs OpenRouter on routing, observability, cost attribution, security, deployment, DX. Why FAGI wins on the self-improving loop in 2026.","href":"/blog/future-agi-vs-openrouter-2026","cat":"Blog"},{"title":"How to Build an LLM Evaluation Framework From Scratch (2026)","desc":"Building an LLM eval framework is a one-week project and a one-year maintenance burden. The eight components, honest cost map, build vs buy guidance.","href":"/blog/build-llm-evaluation-framework-scratch-2026","cat":"Blog"},{"title":"LLM Evaluation Metrics: Everything You Need in 2026","desc":"There aren't 50 LLM eval metrics. Three primitive families and eight rubrics matter in production. 2026 reference with CI gate and per-trace eval cascade.","href":"/blog/llm-evaluation-metrics-everything-you-need-2026","cat":"Blog"},{"title":"Why LLM-as-a-Judge (2026): The Case For, Against, and the Hybrid That Wins","desc":"LLM-as-a-judge is the only eval method that scales to subjective rubrics. The honest case for, against, and the discipline that earns trust.","href":"/blog/why-llm-as-a-judge-2026","cat":"Blog"},{"title":"The Definitive Guide to AI Agent Evaluation (2026)","desc":"2026 working pattern for AI agent evaluation. Six dimensions, six rubrics, 4-D trajectory score, CI gate beats aggregate scoring, loop production needs.","href":"/blog/definitive-guide-ai-agent-evaluation-2026","cat":"Blog"},{"title":"Evaluating LLM Classifiers in 2026: The Eval That Ships","desc":"Evaluating LLM classifiers in 2026: per-class precision-recall, macro vs weighted F1, production-distribution calibration, confusion-matrix debugging.","href":"/blog/evaluating-llm-classifiers-2026","cat":"Blog"},{"title":"LLM Chatbot Evaluation: A Comprehensive Guide (2026)","desc":"Chatbot eval is six stacked problems: intent, retrieval, generation, tool use, multi-turn, safety. One Groundedness score hides what actually ships.","href":"/blog/llm-chatbot-evaluation-comprehensive-guide-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Streaming LLM Responses in 2026","desc":"Five AI gateways for streaming LLM responses in 2026: SSE fidelity, WebSocket and gRPC, mid-stream failover, TTFT observability, CDN.","href":"/blog/best-ai-gateways-streaming-llm-responses-2026","cat":"Blog"},{"title":"Claude Skills Evaluation Deep Dive: The Skill Is the Contract (2026)","desc":"Evaluating Claude Skills in 2026: the skill is a contract, eval the contract. Three rubrics for dispatch, trajectory, integration, on traceAI.","href":"/blog/claude-skills-evaluation-deep-dive-2026","cat":"Blog"},{"title":"Evaluating Claude Sub-Agents: The Dispatch Is the Unit (2026)","desc":"Evaluating Claude sub-agents in 2026: dispatch is the eval unit. Three rubrics, per-handoff scoring in CI, traceAI Task-tool spans, the production loop.","href":"/blog/evaluating-claude-sub-agents-2026","cat":"Blog"},{"title":"Best AI Gateway to Use with Claude Code in 2026","desc":"Five AI gateways scored against Claude Code in 2026: provider breadth, routing, fallback, observability, cost, security, deployment. Opinionated picks.","href":"/blog/best-ai-gateway-to-use-with-claude-code-2026","cat":"Blog"},{"title":"Best 5 Voice AI Simulation Tools for Fintech in 2026","desc":"Five voice AI simulation tools for fintech, KYC, account servicing, fraud callbacks. FFIEC, NYDFS Part 500, FinCEN BSA, CFPB UDAAP, SEC 17a-4.","href":"/blog/best-fintech-voice-ai-simulation-2026","cat":"Blog"},{"title":"Evaluating LLM Data Leakage Prevention (2026)","desc":"Data leakage in LLM systems is four problems, not one. The 2026 method for measuring leak rates across input, output, retrieval, tool-call.","href":"/blog/evaluating-llm-data-leakage-prevention-2026","cat":"Blog"},{"title":"Best 5 Voice AI Simulation Tools for Healthcare in 2026","desc":"Five voice AI simulation tools for healthcare scribes, telehealth triage, medication reminders. HIPAA, HHS OCR, FDA SaMD, ONC HTI-1, BAA.","href":"/blog/best-healthcare-voice-ai-simulation-2026","cat":"Blog"},{"title":"Best 5 Replicate Alternatives in 2026","desc":"Five Replicate alternatives scored on LLM inference depth, catalog breadth, per-token vs per-second economics, custom containers, gateway-in-front pattern.","href":"/blog/best-replicate-alternatives-2026","cat":"Blog"},{"title":"How to Generate Synthetic Data Using LLMs (2026)","desc":"Generate 10x what you need and keep the 10% that survives the filter. Persona, taxonomy, hostile-evolution patterns and the rejection pipeline.","href":"/blog/how-to-generate-synthetic-data-llms-2026","cat":"Blog"},{"title":"Best LLMs of April 2026: Eight Frontier Releases in 30 Days, the Month Trust Broke","desc":"Best LLMs April 2026: compare GPT-5.5, Claude Opus 4.7, DeepSeek V4, Gemma 4, and Qwen after benchmark trust broke and prices compressed fast.","href":"/blog/best-llms-april-2026","cat":"Blog"},{"title":"Best Voice AI Models in April 2026: STT, TTS, and Voice Agent Stack","desc":"Best Voice AI April 2026: compare OpenAI Realtime API, Deepgram, Cartesia, ElevenLabs, Vapi, and Retell for STT, TTS, latency, and voice agents.","href":"/blog/best-voice-ai-april-2026","cat":"Blog"},{"title":"What is Pydantic AI? Type-Safe Agent Framework in 2026","desc":"Pydantic AI is a Python agent framework that brings Pydantic-style validation to LLM tool calls and outputs. Agents, tools, dependency injection, graphs.","href":"/blog/what-is-pydantic-ai-2026","cat":"Blog"},{"title":"Best 5 Evidently AI Alternatives in 2026","desc":"Five Evidently AI alternatives on report-suite portability, LLM-native tracing, guardrails, gateway. What each actually fixes beyond ML-monitoring libs.","href":"/blog/best-evidently-ai-alternatives-2026","cat":"Blog"},{"title":"How to Optimize Vapi Voice Agent Latency in 2026: 12 Techniques + Code","desc":"Optimize Vapi voice agent latency to sub-500ms p95 in 2026. 12 techniques with real Vapi config: streaming STT, partial TTS, prompt caching, regional.","href":"/blog/how-to-optimize-vapi-latency-2026","cat":"Blog"},{"title":"LLM Eval Cost Optimization: 3 Patterns to Cut 80-95% in 2026","desc":"Stack cascade, sampling, and caching to drop LLM judge bills 80 to 95 percent. The math, the code, and the Future AGI surfaces, in one guide.","href":"/blog/llm-eval-cost-optimization-2026","cat":"Blog"},{"title":"Enterprise Controls for All CLI Coding Agents: A 2026 Gateway Field Guide","desc":"A 2026 field guide for governing a mixed CLI coding-agent fleet (Claude Code, Cursor, Codex CLI, Cline, Aider) through one control plane.","href":"/blog/enterprise-controls-cli-coding-agents-gateway-field-guide-2026","cat":"Blog"},{"title":"Evaluating Cohere Rerank in RAG (2026)","desc":"Reranking helps when recall is high but precision is low. It hurts when recall is low. The eval triangle (NDCG@k, recall delta, latency) tells you which.","href":"/blog/evaluating-cohere-rerank-rag-2026","cat":"Blog"},{"title":"LLM Eval Data Warehouse Architecture (2026)","desc":"Trace, eval, cost, outcome, joined on trace_id. The 4-table schema that turns LLM eval dashboards from 'looks bad' to 'fix this prompt.'","href":"/blog/llm-eval-data-warehouse-architecture-2026","cat":"Blog"},{"title":"Agent Rollout Strategies in 2026: The Four-Stage Gate","desc":"Agent rollout is a four-stage gate: shadow, canary, percentage, full. Each stage has a different eval question. Skipping one ships a production incident.","href":"/blog/agent-rollout-strategies-2026","cat":"Blog"},{"title":"LLM Summarization Evaluation: A 2026 Architectural Deep Dive","desc":"Summarization eval is four judge prompts: groundedness, completeness, factuality, conciseness. Each a hardened prompt with a calibration set. 2026 guide.","href":"/blog/llm-summarization-evaluation-deep-dive-2026","cat":"Blog"},{"title":"RAG vs Fine-Tuning: A 2026 Decision Framework","desc":"RAG vs fine-tuning is the wrong question. RAG for facts that change, fine-tune for behavior that doesn't. Three axes, three patterns, the eval.","href":"/blog/rag-vs-fine-tuning-decision-framework-2026","cat":"Blog"},{"title":"Best 5 LangGraph Alternatives in 2026","desc":"Five LangGraph alternatives on framework independence, runtime observability, optimizer + eval native. What each actually fixes when graph-agents plateau.","href":"/blog/best-langgraph-alternatives-2026","cat":"Blog"},{"title":"Evaluating Instructor Structured Outputs (2026)","desc":"Evaluating Instructor structured outputs in 2026: per-field rubrics, cross-field consistency, numeric drift, and traceAI instrumentation.","href":"/blog/evaluating-instructor-structured-outputs-2026","cat":"Blog"},{"title":"Evaluating Streaming LLM Responses in 2026","desc":"Streaming LLM evaluation is four metrics, not one. TTFT, inter-token p99, mid-stream consistency, premature termination. The honest 2026 playbook.","href":"/blog/evaluating-streaming-llm-responses-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Token Budgeting in 2026","desc":"Five AI gateways scored on token budgeting 2026: per-feature allocation, monthly burndown, 75/90/100 alerts, soft vs hard stop, cross-team.","href":"/blog/best-ai-gateways-token-budgeting-2026","cat":"Blog"},{"title":"Best 5 Coval Alternatives in 2026","desc":"Five Coval alternatives on scope beyond voice sim, native gateway, routing, guardrails, optimizer. What each actually fixes after voice-only testing.","href":"/blog/best-coval-alternatives-2026","cat":"Blog"},{"title":"Top 5 AI Gateways to Secure Your Claude Code Rollout in 2026","desc":"Five AI gateways scored on securing a Claude Code rollout: code-egress, secret scanning, prompt injection, audit, IdP, MCP validation.","href":"/blog/top-ai-gateways-secure-claude-code-rollout-2026","cat":"Blog"},{"title":"What is Tokenization in LLMs? BPE, SentencePiece, tiktoken in 2026","desc":"Tokenization explained for 2026 LLMs: BPE, SentencePiece, WordPiece, tiktoken, why tokenizers shape cost, latency, eval scores, and multilingual quality.","href":"/blog/what-is-tokenization-llms-2026","cat":"Blog"},{"title":"Best 5 AIMon Alternatives in 2026","desc":"Five AIMon alternatives on hallucination depth, gateway and routing, optimizer, self-host, migration cost. What each actually fixes beyond hosted scope.","href":"/blog/best-aimon-alternatives-2026","cat":"Blog"},{"title":"PostHog LLM Analytics Alternatives in 2026: 6 Purpose-Built Tools","desc":"FutureAGI closes the self-improving loop for AI product teams; Langfuse, Mixpanel, Amplitude, LangSmith, and Helicone each ship a slice. 2026 picks.","href":"/blog/posthog-llm-analytics-alternatives-2026","cat":"Blog"},{"title":"TrueFoundry Alternatives in 2026: 5 AI Gateway Platforms Compared","desc":"Portkey, Kong AI Gateway, LiteLLM, Helicone, and FutureAGI as TrueFoundry alternatives in 2026. K8s vs hosted, OSS license, and tradeoffs.","href":"/blog/truefoundry-alternatives-2026","cat":"Blog"},{"title":"5 Best AI Answering Services in 2026 (Tested + Ranked)","desc":"Top 5 AI answering services in 2026 ranked on setup speed, integrations, and reliability. Honest tradeoffs plus 2 honorable mentions for SMB owners.","href":"/blog/best-ai-answering-services-2026","cat":"Blog"},{"title":"Best 5 Langtrace Alternatives in 2026","desc":"Five Langtrace alternatives on OTel compatibility, prompt management, hosted pricing. What each actually fixes when OTel-only tracing is no longer enough.","href":"/blog/best-langtrace-alternatives-2026","cat":"Blog"},{"title":"Future AGI vs Bluejay: 2026 Voice Agent Evaluation","desc":"Future AGI vs Bluejay on simulation, native voice observability, eval, inline guardrails, optimizer, pricing, compliance. Honest verdict for voice teams.","href":"/blog/future-agi-vs-bluejay-2026","cat":"Blog"},{"title":"How to Build RAG-Powered Voice AI Agents in 2026","desc":"Build streaming RAG-powered voice agents in 2026. Parallel retrieval, grounded LLM with citations, faithfulness eval, and traceAI instrumented spans.","href":"/blog/how-to-build-rag-powered-voice-ai-agents-2026","cat":"Blog"},{"title":"Automating Lead Qualification with AI Voice Callers in 2026","desc":"Automate outbound lead qualification with AI voice callers in 2026: BANT scoring, objection handling, CRM handoff, pre-launch simulation.","href":"/blog/lead-qualification-ai-voice-callers-2026","cat":"Blog"},{"title":"Multi-Agent Voice Systems in 2026: State Transitions, Hand-offs, Eval Boundaries","desc":"How to architect multi-agent voice systems in 2026: state transitions, hand-off prompt design, per-agent vs e2e evals, latency budgets, attribution.","href":"/blog/multi-agent-voice-systems-2026","cat":"Blog"},{"title":"Three-Layer Voice AI Testing: Regression, Adversarial, Production-Derived","desc":"The 2026 voice testing pattern: regression on golden conversations, adversarial red-team personas, production-derived replays. Engineering build guide.","href":"/blog/three-layer-voice-testing-2026","cat":"Blog"},{"title":"The Voice Agent Eval Rubric Library: 14 Rubrics for 2026","desc":"The 14 voice-specific eval rubrics worth running in 2026. Named templates, when to use each, what bad looks like, code samples, how they layer.","href":"/blog/voice-agent-eval-rubric-library-2026","cat":"Blog"},{"title":"Voice AI Observability for Pipecat: A 2026 Implementation Guide","desc":"Implement voice observability for Pipecat with traceAI-pipecat: install, register, enable HTTP attribute mapping, attach audio + multi-turn eval rubrics.","href":"/blog/voice-ai-observability-pipecat-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for MCP Tool-Level Observability with Codex CLI in 2026","desc":"Five AI gateways scored for MCP tool-level observability with Codex CLI: per-tool latency, success rate, argument validation, MCP auth.","href":"/blog/best-ai-gateways-mcp-tool-observability-codex-cli-2026","cat":"Blog"},{"title":"Best CX AI Evaluation Platforms in 2026: 5 Picks","desc":"Five CX AI evaluation platforms scored on CustomerAgent rubrics, paired Containment and False-Resolution KPIs, and Zendesk/Intercom span attribution.","href":"/blog/best-cx-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"Evaluating vLLM Self-Hosted LLMs in 2026","desc":"How to evaluate a vLLM self-hosted LLM in 2026: catch continuous-batching jitter, KV-cache eviction, and AWQ/GPTQ/FP8 drift before prod.","href":"/blog/evaluating-vllm-self-hosted-llm-2026","cat":"Blog"},{"title":"LLM Eval Monitoring Dashboards: The Four Panels That Drive Action","desc":"Eval dashboards are not Grafana for LLM logs. Four panels that actually drive action: rubric trends, per-route delta, Error Feed top-N, cost/resolved.","href":"/blog/llm-eval-monitoring-dashboards-2026","cat":"Blog"},{"title":"Your Agent Passes Evals and Fails in Production. Here's Why. (2026)","desc":"Your eval set is a snapshot, production is a river. Six drift modes that age eval sets, and the trace-as-eval loop that closes the gap.","href":"/blog/agent-passes-evals-fails-production-2026","cat":"Blog"},{"title":"Best 5 Lunary Alternatives in 2026","desc":"Five Lunary alternatives on community size, hosted pricing, gateway and optimizer depth. What each actually fixes once you outgrow lightweight tracing.","href":"/blog/best-lunary-alternatives-2026","cat":"Blog"},{"title":"Evaluating LLM Citation & Attribution (2026)","desc":"Citation eval is three rubrics: did the model emit a citation, does it resolve, does the source actually contain the claim. 2026 methodology with code.","href":"/blog/evaluating-llm-citation-attribution-2026","cat":"Blog"},{"title":"Evaluating LLM Content Moderation (2026)","desc":"Evaluate LLM content moderation in 2026: the 2x2 matrix per category, adversarial and benign test sets, threshold tuning, and the monthly drift loop.","href":"/blog/evaluating-llm-content-moderation-2026","cat":"Blog"},{"title":"Best 5 Anyscale Alternatives for LLM Workloads in 2026","desc":"Five Anyscale alternatives on LLM-native surface, inference cost at scale, gateway, optimizer. What each actually fixes for LLM-first vs Ray-first work.","href":"/blog/best-anyscale-llm-alternatives-2026","cat":"Blog"},{"title":"Evaluating DeepSeek Models in 2026","desc":"Evaluate DeepSeek V3, R1, and V4 for production: capability-shape benchmarks, paired comparison on YOUR English data, safety regression, residency gate.","href":"/blog/evaluating-deepseek-models-2026","cat":"Blog"},{"title":"LLM Eval Time-to-Value: A 2026 Milestone-by-Milestone Rollout","desc":"Practical time-to-value plan for your LLM eval stack: day 1-3 smoke set, week 1 PR gate, month 1 incident clustering, quarter 1 budget chargeback.","href":"/blog/llm-eval-time-to-value-2026","cat":"Blog"},{"title":"LLM-Judge Prompt Engineering: The 2026 Engineering Guide","desc":"Judge prompts are eval-time programs. Five elements decide signal vs noise: criterion, calibrated examples, randomization, schema, self-consistency.","href":"/blog/llm-judge-prompt-engineering-guide-2026","cat":"Blog"},{"title":"How to Build and Evaluate a Customer Support Chatbot in 2026","desc":"Customer support eval in 2026: escalation taxonomy first, clause-level retrieval, tool-call correctness on Zendesk and Intercom, paired Containment.","href":"/blog/customer-support-chatbot-build-evaluate-2026","cat":"Blog"},{"title":"Evaluating LLM Agent Handoffs (2026)","desc":"Evaluating LLM agent handoffs in 2026: the handoff is the cross-framework eval unit. Four rubrics, per-handoff spans, CI gates, and Error Feed clustering.","href":"/blog/evaluating-llm-agent-handoffs-2026","cat":"Blog"},{"title":"Evaluating Tool-Calling Agents in 2026: The Four-Layer Eval Stack","desc":"Tool-calling eval is four problems stacked: tool selection, argument extraction, result utilization, error recovery. Most posts grade only the first.","href":"/blog/evaluating-tool-calling-agents-2026","cat":"Blog"},{"title":"The Ultimate Guide to LLM Guardrails (2026)","desc":"Senior-engineer guide to LLM guardrails: placement, 9 open-weight + 4 API backends, latency budgets, ensembles, precision/recall that actually catches.","href":"/blog/ultimate-guide-llm-guardrails-2026","cat":"Blog"},{"title":"Autoresearch LLM Test Generation: 2026 Loop That Actually Works","desc":"Autoresearch LLM test generation as a hostile-judge loop: persona x scenario x adversarial evolution, cross-family judge, keep top 10 percent.","href":"/blog/autoresearch-llm-test-generation-2026","cat":"Blog"},{"title":"Best LLM Prompt Playgrounds in 2026: 7 Tools Compared","desc":"FutureAGI, Langfuse, OpenAI, Anthropic, PromptLayer, Helicone, and Vercel AI Playground for LLM prompt iteration in 2026. Diff, version, score, deploy.","href":"/blog/best-llm-prompt-playground-tools-2026","cat":"Blog"},{"title":"Evaluating smolagents in 2026: Code-as-Action Eval","desc":"smolagents' CodeAgent makes the plan AS code, so the eval changes shape: code synthesis correctness, sandbox safety, and result-interpretation fidelity.","href":"/blog/evaluating-smolagents-2026","cat":"Blog"},{"title":"OpenAI Frontier vs Claude Cowork: Enterprise Agents Compared (2026)","desc":"OpenAI Frontier vs Claude Cowork 2026 head-to-head: agent execution, governance, security, pricing, and the eval layer every CTO needs on top of both.","href":"/blog/openai-frontier-vs-claude-cowork-enterprise-comparison","cat":"Blog"},{"title":"Best 5 AI Gateways for Devin-Style Autonomous Coding Agents in 2026","desc":"Five AI gateways scored on Devin-style autonomous coding agent workloads in 2026: long-session traces, per-task spend caps, scoring, failure clusters.","href":"/blog/best-ai-gateways-devin-autonomous-coding-agents-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for LLM Failover and Fallback in 2026","desc":"Five AI gateways for LLM failover and fallback in 2026 scored on health-detection latency, streaming continuity, idempotency, MTTR, fallback-route quality.","href":"/blog/best-ai-gateways-llm-failover-fallback-2026","cat":"Blog"},{"title":"Best 5 PromptLayer Alternatives in 2026","desc":"Five PromptLayer alternatives on multi-provider routing, observability depth, optimizer, lower-tier RBAC, SDK breadth. What each actually fixes.","href":"/blog/best-promptlayer-alternatives-2026","cat":"Blog"},{"title":"Evaluating LlamaIndex RAG Applications in 2026","desc":"LlamaIndex RAG eval is not generic RAG eval. Four layers, four rubrics: retriever, query-pipeline, agent tool calls, and the traceAI bridge to production.","href":"/blog/evaluating-llamaindex-rag-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Manufacturing in 2026: OT/IT Data Boundaries, Latency Budgets, and Edge Routing","desc":"Five AI gateways for manufacturing 2026 scored on IEC 62443 zones, ISA-95 OT/IT boundary, NIST CSF 2.0, ITAR residency, sub-500ms latency, offline-first.","href":"/blog/best-ai-gateways-manufacturing-2026","cat":"Blog"},{"title":"Best 5 vLLM Alternatives for Self-Hosted Inference in 2026","desc":"Five vLLM alternatives scored on throughput, hardware coverage, quantization, structured outputs, ops burden, plus the platform layer that augments them.","href":"/blog/best-vllm-self-hosted-inference-alternatives-2026","cat":"Blog"},{"title":"How to Evaluate Voice AI Agents End-to-End: A 2026 Methodology","desc":"Step-by-step 2026 methodology to evaluate voice AI agents end-to-end: trace, score, cluster, optimize, redeploy. With real rubrics, code, a closed loop.","href":"/blog/how-to-evaluate-voice-ai-agents-end-to-end-2026","cat":"Blog"},{"title":"Voice Cloning Safety and Brand Voice Management for Production AI in 2026","desc":"Manage voice cloning safety and brand voice for production AI in 2026 with consent capture, watermarking, voice-print policy, and Future AGI Protect.","href":"/blog/voice-cloning-safety-brand-voice-2026","cat":"Blog"},{"title":"Claude Fortress Analysis (Defender's View): Why Model Safety Isn't Enough","desc":"A defender's analysis of why Claude is the hardest frontier model to break in 2026, where Constitutional AI earns it, where the fortress cracks.","href":"/blog/claude-fortress-defender-analysis-2026","cat":"Blog"},{"title":"External Evaluation Pipelines for LLM Apps in 2026","desc":"External eval pipelines need four properties: async, idempotent, observable, recoverable. Working blueprint with FAGI distributed runners and real code.","href":"/blog/external-evaluation-pipelines-llm-apps-2026","cat":"Blog"},{"title":"How to Build and Evaluate a Medical Chatbot in 2026","desc":"Medical chatbot eval in 2026: refusal as primary metric, four-tier taxonomy, citation enforcement on every claim, PHI guardrails, clinical-pilot list.","href":"/blog/medical-chatbot-build-evaluate-2026","cat":"Blog"},{"title":"Red Teaming LLMs: A Step-by-Step Guide (2026)","desc":"Red-teaming an LLM is three loops: probe, classify, triage. A 2026 playbook that wires PyRIT and garak into a continuous compounding CI gate.","href":"/blog/red-teaming-llms-step-by-step-2026","cat":"Blog"},{"title":"Best AI Agent Failure Detection Tools in 2026: 6 Compared","desc":"Six AI agent failure detection tools for ML and SRE teams 2026: eval-on-every-span, auto-clustering, runtime guards, alert routing, what actually pages.","href":"/blog/best-ai-agent-failure-detection-tools-2026","cat":"Blog"},{"title":"Best 5 Comet ML Alternatives in 2026","desc":"Five Comet ML alternatives on LLM-native tracing, OpenInference/OTel, gateway and optimizer. What each actually fixes when workload moves to agent traces.","href":"/blog/best-comet-ml-alternatives-2026","cat":"Blog"},{"title":"Best LLMOps Platforms in 2026: 7 End-to-End Stacks Compared","desc":"FutureAGI, Langfuse, MLflow, W&B Weave, Comet, Braintrust, LangSmith for LLMOps in 2026. Pricing, OSS license, and what each platform won't do end-to-end.","href":"/blog/best-llmops-platforms-2026","cat":"Blog"},{"title":"Best 5 Self-Hosted AI Gateways in 2026","desc":"Five self-hosted AI gateways scored on Kubernetes and Docker support, HA topology, self-observability, upgrade path, license, footprint, ops burden.","href":"/blog/best-self-hosted-ai-gateways-2026","cat":"Blog"},{"title":"Best 5 Labelbox Alternatives for LLM Workflows in 2026","desc":"Five Labelbox alternatives on eval-dataset portability, runtime traces, inline guardrails, optimizer. What each actually fixes when annotation-first lags.","href":"/blog/best-labelbox-llm-alternatives-2026","cat":"Blog"},{"title":"Evaluating Voice AI Agents in 2026: The Methodology","desc":"Voice agent eval is end-task scoring plus pipeline-stage attribution plus conversation coherence. WER scores the ASR component, not the agent.","href":"/blog/evaluating-voice-ai-agents-2026","cat":"Blog"},{"title":"Future AGI vs Cloudflare AI Gateway in 2026: Self-Improving Runtime vs Edge Cache","desc":"Future AGI vs Cloudflare AI Gateway scored on routing, observability, cost attribution, security, deployment, DX. The honest verdict and pricing snapshot.","href":"/blog/future-agi-vs-cloudflare-ai-gateway-2026","cat":"Blog"},{"title":"How an MCP Gateway Cuts Token Costs in Claude Code and Codex CLI in 2026","desc":"A 2026 architecture essay on why MCP blows up coding-agent token bills in Claude Code and Codex CLI, and five mechanisms that compress cost.","href":"/blog/how-mcp-gateway-cuts-token-costs-claude-code-codex-cli-2026","cat":"Blog"},{"title":"Fine-Tuning Pipeline Evaluation: A 2026 Deep Dive","desc":"Pipeline-level eval for fine-tuning in 2026: four stages, four checks. Data quality, held-out plus drift, paired vs base, production canary.","href":"/blog/fine-tuning-pipeline-evaluation-2026","cat":"Blog"},{"title":"Multilingual LLM Evaluation: A 2026 Playbook for Non-English","desc":"Ship LLM eval that holds up outside English: 7 multilingual challenges, 5-step rollout, classifier ensembles, and the FAGI grounding loop.","href":"/blog/llm-eval-multilingual-non-english-2026","cat":"Blog"},{"title":"The 2026 LLM Evaluation Playbook","desc":"The pillar playbook for LLM evaluation in 2026: dataset, metrics, judge, CI gate, production observation, closed loop from failing trace to regression.","href":"/blog/llm-evaluation-playbook-2026","cat":"Blog"},{"title":"Top LLM Evaluators for Testing LLMs at Scale (2026)","desc":"Scaling LLM tests is three primitives: distributed runners, classifier cascade, per-route sampling. Six evaluators ranked by burst survival.","href":"/blog/top-llm-evaluators-testing-at-scale-2026","cat":"Blog"},{"title":"Best 5 Arize Phoenix Alternatives in 2026","desc":"Five Arize Phoenix alternatives scored on prompt management, gateway integration, evaluation maturity, fixes once OSS-only tracing stops being enough.","href":"/blog/best-arize-phoenix-alternatives-2026","cat":"Blog"},{"title":"Future AGI vs Vercel AI Gateway in 2026: Self-Improving Runtime vs Framework-Native Proxy","desc":"Future AGI vs Vercel AI Gateway scored on routing, observability, cost attribution, security, deployment, DX. Honest verdict and pricing snapshot for 2026.","href":"/blog/future-agi-vs-vercel-ai-gateway-2026","cat":"Blog"},{"title":"How to Reduce Claude Code Token Costs by Up to 90 Percent in 2026","desc":"Cut Claude Code token spend with 5 stackable levers: cache_control, MCP-tool compilation, semantic caching, model right-sizing, pruning. Honest 90% read.","href":"/blog/how-to-reduce-claude-code-token-costs-90-percent-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for SaaS Platforms in 2026: Multi-Tenant Routing, Cost Attribution, and Usage-Based Billing","desc":"Five AI gateways for B2B SaaS in 2026: multi-tenant isolation, BYO-key segmentation, per-tenant cost for usage billing, GDPR DPA, fair-share rate limits.","href":"/blog/best-ai-gateways-saas-2026","cat":"Blog"},{"title":"Best 5 Maxim Bifrost Alternatives in 2026","desc":"Five Maxim Bifrost alternatives on bundle independence, OSS instrumentation, observability, community. What each actually fixes when Maxim stops fitting.","href":"/blog/best-maxim-bifrost-alternatives-2026","cat":"Blog"},{"title":"Opik Alternatives in 2026: 6 LLM Eval and Observability Tools","desc":"FutureAGI, Langfuse, Phoenix, Braintrust, LangSmith, and DeepEval as Comet Opik alternatives in 2026. Pricing, OSS license, judge metrics, and tradeoffs.","href":"/blog/comet-opik-alternatives-2026","cat":"Blog"},{"title":"LangChain Callback Tracing Best Practices 2026: Spans, Cardinality","desc":"LangChain callback tracing best practices in 2026: handler design, async support, cardinality, span hierarchy, OTel integration, when to skip callbacks.","href":"/blog/langchain-callback-tracing-best-practices-2026","cat":"Blog"},{"title":"12 Metrics for AI Conversation Monitoring in 2026","desc":"Twelve metrics across five axes for monitoring conversational agents in production: coherence, resolution, safety, cost, adaptation. Wiring included.","href":"/blog/12-metrics-ai-conversation-monitoring-2026","cat":"Blog"},{"title":"AI Call Center QA Software: 8 Tools Compared (2026)","desc":"AI call center QA software 2026 splits into transcript scoring and live coaching. Eight tools on rubric depth, real-time latency, deployment, honest fit.","href":"/blog/ai-call-center-qa-software-2026","cat":"Blog"},{"title":"Audio Caching for Voice AI: 2026 Latency Reduction Guide","desc":"Audio caching is a quarter of the voice AI latency story. The full guide: semantic cache at the orchestrator, TTS prefix cache, streaming, p95 honestly.","href":"/blog/audio-caching-latency-reduction-voice-ai-2026","cat":"Blog"},{"title":"Cascaded Voice AI vs Speech-to-Speech: The 2026 Architecture Decision","desc":"Cascaded voice AI vs speech-to-speech in 2026: latency, eval depth, debug cost, model flexibility, and the architecture decision every voice team faces.","href":"/blog/cascaded-voice-ai-vs-speech-to-speech-2026","cat":"Blog"},{"title":"ElevenLabs vs Cartesia: 2026 Streaming TTS Deep Comparison","desc":"ElevenLabs vs Cartesia in 2026: streaming TTFA latency, voice realism, cloning, multilingual coverage, SSML, pricing, same-rubric evaluation guide.","href":"/blog/elevenlabs-vs-cartesia-tts-2026","cat":"Blog"},{"title":"Future AGI vs Coval in 2026: Closed-Loop Voice Platform vs Focused Simulation","desc":"Future AGI vs Coval on simulation, native voice observability, eval, inline guardrails, optimization, pricing, compliance. Honest verdict, May 2026.","href":"/blog/future-agi-vs-coval-2026","cat":"Blog"},{"title":"How to Use Voice Agent Analytics to Improve CSAT in 2026","desc":"Use voice agent analytics to lift CSAT in 2026. Instrument calls with traceAI, score with CSAT-proxy rubrics, cluster failures with Error Feed, fix.","href":"/blog/how-to-improve-voice-agent-csat-with-analytics-2026","cat":"Blog"},{"title":"Voice Agent Deployment Patterns: Cloud, BYOC, and On-Prem in 2026","desc":"Three voice agent deployment patterns compared in 2026. Cloud (managed hosted), BYOC inside customer VPC, and air-gapped on-prem with concrete tradeoffs.","href":"/blog/voice-agent-deployment-patterns-cloud-byoc-2026","cat":"Blog"},{"title":"7 Voice Agent ASR Failure Modes in Production (and How to Catch Them)","desc":"The 7 ASR failure modes that break voice agents in production: detection patterns via spans, rubrics, Error Feed clusters, mitigation plays.","href":"/blog/voice-agent-asr-failure-modes-2026","cat":"Blog"},{"title":"Logging and Analytics Architecture for Voice Agents in 2026","desc":"Design the data plane for voice agents in 2026: spans, OTLP, eval, dashboards, alerts, retention, and GDPR/HIPAA tradeoffs across the full architecture.","href":"/blog/voice-agent-logging-analytics-architecture-2026","cat":"Blog"},{"title":"Voice AI Drop-Off Rate: The Metric That Predicts Hang-Up Risk","desc":"Drop-off rate beats CSAT as leading indicator. Tag traces, score with conversation_resolution and task_completion, pinpoint the turn that caused hang-up.","href":"/blog/voice-ai-drop-off-rate-metric-2026","cat":"Blog"},{"title":"Voice AI for Insurance: Claims Intake, Underwriting Calls, and Compliance in 2026","desc":"How to deploy voice AI across insurance workflows in 2026. FNOL intake, claims status, underwriting Q&A, dispatch, renewals, fraud triage, compliance.","href":"/blog/voice-ai-insurance-claims-underwriting-2026","cat":"Blog"},{"title":"Voice AI for Legal: Discovery Intake, Client Onboarding, and Compliance in 2026","desc":"Deploy voice AI in legal workflows in 2026: client intake, discovery interviews, contract Q&A, deposition prep, status updates, and the compliance posture.","href":"/blog/voice-ai-legal-discovery-intake-2026","cat":"Blog"},{"title":"Voice AI Observability for LiveKit Agents: A 2026 Guide","desc":"Implement voice AI observability for LiveKit Agents: native FAGI dashboard via Assistant ID plus traceai-livekit pip package for code-driven span tracing.","href":"/blog/voice-ai-observability-livekit-2026","cat":"Blog"},{"title":"Voice AI Load Testing: Simulating 10,000+ Concurrent Calls in 2026","desc":"Load test voice AI at 10,000+ concurrent calls in 2026: spawn parallel personas, score under load, find latency degradation and eval drift before ship.","href":"/blog/voice-load-testing-simulating-10000-calls-2026","cat":"Blog"},{"title":"Why WER Isn't Enough for Voice Agents: 2026 Beyond-WER Metrics","desc":"WER measures word accuracy but misses what voice agents break on. Intent preservation, entity F1, timing, task-completion correlation are 2026 metrics.","href":"/blog/wer-voice-agents-beyond-2026","cat":"Blog"},{"title":"Best 5 Janus AI Alternatives in 2026","desc":"Five Janus AI alternatives on integrated observability, eval, optimizer, gateway, routing, self-host. What each actually fixes vs a hosted agent-builder.","href":"/blog/best-janus-ai-alternatives-2026","cat":"Blog"},{"title":"Multimodal LLM Tracing in 2026: The Methodology That Actually Works","desc":"Multimodal LLM tracing for Gemini Vision, GPT-5 Vision, Claude Vision. Modality attribution, per-modality cost, PII at boundary, traceAI schema.","href":"/blog/multi-modal-llm-tracing-2026","cat":"Blog"},{"title":"How to Build (and Evaluate) a PDF QA Chatbot in 2026","desc":"A PDF QA chatbot is a retrieval problem, not generation. Parse, chunk, hybrid retrieve, cite, evaluate retrieval, bridge to OTel spans.","href":"/blog/pdf-qa-chatbot-build-evaluate-2026","cat":"Blog"},{"title":"Phoenix Alternatives in 2026: 6 LLM Tracing and Eval Platforms","desc":"FutureAGI, Langfuse, LangSmith, Helicone, Braintrust, and W&B Weave as Arize Phoenix alternatives in 2026. Pricing, OSS license, OTel coverage, tradeoffs.","href":"/blog/phoenix-alternatives-2026","cat":"Blog"},{"title":"Best 5 Jaxon AI Alternatives in 2026","desc":"Five Jaxon AI alternatives on synthetic-data depth, gateway, observability, optimizer, languages. What each actually fixes beyond synth-data-only tooling.","href":"/blog/best-jaxon-ai-alternatives-2026","cat":"Blog"},{"title":"Evaluating Agent Memory Systems in 2026: Four Dimensions Most Reports Miss","desc":"Evaluating agent memory is four problems, not one: recall, freshness, contradiction handling, forgetting. A 2026 framework for Mem0, Zep, Letta, LangMem.","href":"/blog/evaluating-agent-memory-systems-2026","cat":"Blog"},{"title":"Evaluating Search-Augmented Agents in 2026","desc":"Generic RAG eval misses what kills search agents: bad queries, stale sources, monoculture, and broken cites. A four-axis rubric you can ship this week.","href":"/blog/evaluating-search-augmented-agents-2026","cat":"Blog"},{"title":"LangGraph Agent Evaluation: A 2026 Deep Tutorial","desc":"LangGraph eval is graph-level, not message-level. Score state transitions: node-input, node-output, edge-routing, and checkpoint replay determinism.","href":"/blog/langgraph-agent-evaluation-2026","cat":"Blog"},{"title":"Best 5 AI Observability Tools for Cybersecurity in 2026","desc":"Cybersecurity AI observability in 2026: five platforms scored on per-request span, SIEM export, prompt-injection detection at the trace layer. FAGI, DD.","href":"/blog/best-cybersecurity-ai-observability-2026","cat":"Blog"},{"title":"Eval SDK vs Eval Platform vs Build: The 2026 Decision Framework","desc":"Build vs buy for LLM evaluation 2026. SDK vs hosted platform tradeoffs across seven axes, cost math, the hybrid pattern most production teams should run.","href":"/blog/eval-sdk-vs-eval-platform-build-vs-buy-2026","cat":"Blog"},{"title":"LLM Eval vs Fine-Tuning: When to Do What in 2026","desc":"When eval-driven prompt optimization is enough vs fine-tuning, the seven decision axes, and the five-step path that ships most teams without retrain.","href":"/blog/llm-eval-vs-fine-tuning-when-to-do-what-2026","cat":"Blog"},{"title":"LLM Testing in 2026: Methods and Strategies","desc":"The 2026 LLM testing pyramid: deterministic unit checks, regression eval, red-team and canary at the top, wired into pytest and CI.","href":"/blog/llm-testing-2026-methods-strategies","cat":"Blog"},{"title":"Best AI Gateway for OpenHands and SWE-Agent Autonomous Workflows in 2026","desc":"Five AI gateways scored on OpenHands and SWE-Agent workflows 2026: per-issue cost caps, trajectory observability, model routing, loop safety.","href":"/blog/best-ai-gateway-openhands-swe-agent-autonomous-workflows-2026","cat":"Blog"},{"title":"Best MCP Gateway for Claude Code to Cut Token Costs by 50 Percent in 2026","desc":"MCP gateway in front of Claude Code cuts input-token spend 50% in 2026: compiled tools, semantic caching, registration, scored across 5 real gateways.","href":"/blog/best-mcp-gateway-claude-code-cut-token-costs-50-percent-2026","cat":"Blog"},{"title":"The Comprehensive Guide to LLM Security (2026)","desc":"LLM security is four layers: input, output, retrieval, tool-call. Defenders that cover all four ship; input-only defenders lose to anything.","href":"/blog/comprehensive-guide-llm-security-2026","cat":"Blog"},{"title":"Future AGI vs Helicone in 2026: Self-Improving Runtime vs Lightweight Observability","desc":"Future AGI vs Helicone scored on instrumentation, observability depth, evaluation, optimization, deployment, DX. Honest verdict and Mintlify posture.","href":"/blog/future-agi-vs-helicone-2026","cat":"Blog"},{"title":"AI Safety Engineering in 2026: CI Guardrails, Drift, and Monitoring","desc":"How engineering teams ship safe AI in 2026. CI/CD guardrails, drift detection, adversarial robustness, monitoring. Future AGI Protect + Guardrails as #1.","href":"/blog/ai-safety-engineering-teams-production-workflow","cat":"Blog"},{"title":"Best Prompt Testing Frameworks in 2026: 7 Compared","desc":"Promptfoo, FutureAGI, Braintrust, LangSmith, Inspect AI, MLflow, OpenPipe for prompt testing in 2026. Compared on regression, red-team, A/B, and CI gating.","href":"/blog/best-prompt-testing-frameworks-2026","cat":"Blog"},{"title":"LLM Latency Tail Evaluation: p99 Methodology for 2026","desc":"P50 lies, p99 tells truth. The 2026 method for measuring throttling, retries, and cache misses per component with traceAI and ACC.","href":"/blog/llm-latency-tail-evaluation-2026","cat":"Blog"},{"title":"Using OpenAI Codex CLI with Multiple Model Providers in 2026: A Gateway Setup Guide","desc":"Walkthrough for pointing OpenAI Codex CLI at Anthropic, Gemini, Mistral, OSS models through an AI gateway in 2026, with 5 gateway picks scored.","href":"/blog/openai-codex-cli-multiple-model-providers-gateway-setup-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Routing Claude Code Requests in Production in 2026","desc":"Five AI gateways scored on routing Claude Code requests in production: policy expressiveness, per-region routing, failover, P99 overhead, observability.","href":"/blog/best-ai-gateways-routing-claude-code-production-2026","cat":"Blog"},{"title":"Best 5 Kong AI Gateway Alternatives in 2026","desc":"Five Kong AI Gateway alternatives scored on AI Proxy plugin portability, observability depth, native eval and optimizer surfaces, pricing above $1.5K/mo.","href":"/blog/best-kong-ai-gateway-alternatives-2026","cat":"Blog"},{"title":"Best 5 LangChain Alternatives for Production LLM Stack in 2026","desc":"Five LangChain alternatives on debug surface, breaking-change cadence, native gateway, optimizer, deps. What each actually fixes when chains stop scaling.","href":"/blog/best-langchain-production-llm-alternatives-2026","cat":"Blog"},{"title":"Evaluating DSPy Pipelines in 2026: Signature-Level Eval","desc":"Evaluating DSPy pipelines in 2026: why the compile metric isn't your production rubric, and how to eval the Signature instead of the program.","href":"/blog/evaluating-dspy-pipelines-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Insurance in 2026: Claims, Underwriting, and Audit Logs","desc":"Five AI gateways for P&C, Life, and Health insurance in 2026 scored on NAIC Model Bulletin, NYDFS CL No. 7, EU AI Act Annex III, HIPAA, CO Reg 10-1-1.","href":"/blog/best-ai-gateways-insurance-2026","cat":"Blog"},{"title":"Best 5 Collinear AI Alternatives in 2026","desc":"Five Collinear AI alternatives on eval breadth, gateway and runtime, optimizer, languages. What each actually fixes after alignment-only stacks plateau.","href":"/blog/best-collinear-ai-alternatives-2026","cat":"Blog"},{"title":"Best 5 Lakera Guard Alternatives in 2026","desc":"Five Lakera Guard alternatives scored on inline guardrail latency, native gateway and eval surfaces, self-host posture, sub-enterprise pricing.","href":"/blog/best-lakera-guard-alternatives-2026","cat":"Blog"},{"title":"LLM Eval vs Traditional QA: A 2026 Bridge for QA Teams","desc":"Traditional QA asserts pass/fail. LLM eval grades against a rubric. The pyramid, the golden set, the CI gate carry, the assertion library you replace.","href":"/blog/llm-eval-vs-traditional-qa-2026","cat":"Blog"},{"title":"Best AI Gateways for A/B Testing LLM Models and Prompts in 2026","desc":"Five AI gateways for A/B testing LLM models and prompts in 2026, scored on shadow traffic, sample-size enforcement, outcome-attached gateway-hop scoring.","href":"/blog/best-ai-gateways-ab-testing-llm-models-prompts-2026","cat":"Blog"},{"title":"How to Build (and Evaluate) a Contract Review RAG Agent in 2026","desc":"Contract review RAG in 2026: clause-level retrieval, citation enforcement, the eval suite in-house counsel signs off on, LangGraph wiring to OTel traces.","href":"/blog/contract-review-rag-build-evaluate-2026","cat":"Blog"},{"title":"Evaluating Cline and Cursor Coding Agents in 2026: A Tutorial","desc":"Cline and Cursor solve different problems. Cline is a VSCode extension with MCP and BYOK. Cursor is a full IDE with Composer. Eval them differently.","href":"/blog/evaluating-cline-cursor-coding-agents-2026","cat":"Blog"},{"title":"LLM Observability Platform Buyer's Guide 2026","desc":"2026 buyer guide for LLM observability platforms: 10 criteria, 7 vendor categories, the 5-question vendor interview, an honest and calibrated ranking.","href":"/blog/llm-observability-platform-buyers-guide-2026","cat":"Blog"},{"title":"Best LLMs of March 2026: When Open-Weight Caught Closed-Source on Coding","desc":"Best LLMs March 2026: compare Gemini 3.1 Pro, Claude Opus 4.6, Mistral Small 4, and Qwen for coding, cost, multimodal, and open-weight picks.","href":"/blog/best-llms-march-2026","cat":"Blog"},{"title":"Best Voice AI Models in March 2026: STT, TTS, and Voice Agent Stack","desc":"Best Voice AI March 2026: Deepgram, Cartesia, ElevenLabs, Vapi, Retell across STT, TTS, latency, and voice agents.","href":"/blog/best-voice-ai-march-2026","cat":"Blog"},{"title":"CI/CD LLM Eval with GitHub Actions: 2026 Workflow","desc":"Cheap, fast, statistically significant LLM eval gates in GitHub Actions: classifier cascade, fi CLI exit codes, Welch's t-test, auto-rollback.","href":"/blog/ci-cd-llm-eval-github-actions-2026","cat":"Blog"},{"title":"Evaluating Fine-Tuned LLMs: A 2026 Playbook","desc":"Fine-tune eval in 2026 without the theatre: four-set gap, paired arena against base, bootstrap CI math, CI gate in code, production canary on spans.","href":"/blog/evaluating-fine-tuned-llms-2026","cat":"Blog"},{"title":"Best 7 AI Gateways for Multi-Model Routing in 2026","desc":"Seven AI gateways for multi-model LLM routing in 2026, ranked on the Future AGI Gateway Scorecard. Covers 15 routing strategies plus the trust cohort.","href":"/blog/best-ai-gateways-model-routing","cat":"Blog"},{"title":"G-Eval (2026): The Definitive Guide for Production LLM Teams","desc":"G-Eval 2026: what the paper actually shipped, where the method breaks in production, four biases that wreck a rubric judge, how to harden for real traffic.","href":"/blog/g-eval-definitive-guide-2026","cat":"Blog"},{"title":"How to Optimize Pipecat Voice Agent Latency in 2026: 12 Techniques + Code","desc":"Cut Pipecat voice agent latency to sub-500ms p95 in 2026. 12 techniques with real pipeline code: streaming STT, partial TTS, prefix caching, routing.","href":"/blog/how-to-optimize-pipecat-latency-2026","cat":"Blog"},{"title":"LLM Eval Feedback Loop Design: A 2026 Engineering Guide","desc":"How to design an LLM evaluation feedback loop that compounds: capture, join, calibrate, promote to dataset, gate in CI, and optimize. The six-stage shape and the dataset-promotion step most teams skip.","href":"/blog/llm-eval-feedback-loop-design-2026","cat":"Blog"},{"title":"LLM Eval Team Scaling Guide 2026: From 5 to 500 Engineers","desc":"How the LLM eval function grows non-linearly from 5 to 500 engineers: five stages, four hand-off inflection points, anti-patterns, FAGI primitives.","href":"/blog/llm-eval-team-size-scaling-guide-2026","cat":"Blog"},{"title":"LLM Eval vs Classical ML Eval: A 2026 Bridge for ML Teams","desc":"Classical ML eval is closed-form, LLM eval is open-form. The discipline that carries, metrics that break, mapping sklearn to an LLM eval suite.","href":"/blog/llm-eval-vs-classical-ml-eval-2026","cat":"Blog"},{"title":"RAG vs CAG: Choosing Cache-Augmented Generation in 2026","desc":"RAG vs Cache-Augmented Generation in 2026: 7 axes for choosing, the hybrid router pattern most teams ship, how to eval both paths with traceAI and FAGI.","href":"/blog/rag-vs-cag-cache-augmented-generation-2026","cat":"Blog"},{"title":"Ragas vs Future AGI in 2026: RAG-Only Library vs Eval-Stack Package","desc":"Ragas vs Future AGI compared honestly on RAG metric coverage, cost economics, CI fit, runtime guardrails, observability, and the closed loop. Where each one wins, where they tie, and how to compose.","href":"/blog/ragas-vs-future-agi-2026","cat":"Blog"},{"title":"Evaluating Mistral Agents in 2026: The Surprises","desc":"Evaluating Mistral agents: the tool-call schema parsing gap, system-prompt adherence vs OpenAI, EU data-residency verification, and Codestral safety gates.","href":"/blog/evaluating-mistral-agents-2026","cat":"Blog"},{"title":"LLM Evaluation Best Practices Checklist for 2026","desc":"7-item LLM eval best practices checklist that actually ships: dataset, judge calibration, deterministic floor, CI gate, stats, observability, closed loop.","href":"/blog/llm-evaluation-best-practices-checklist-2026","cat":"Blog"},{"title":"The 2026 LLM Incident Response Playbook","desc":"LLM incidents need a different playbook than API incidents. Six steps, four incident classes, postmortem becomes golden-set entry to loop.","href":"/blog/llm-incident-response-playbook-2026","cat":"Blog"},{"title":"Prompt Versioning and Lifecycle Management in 2026","desc":"Prompt versioning is git for prompts plus eval-gated promotion plus production rollback. Three-stage lifecycle (draft, gated, deprecation) with FAGI tools.","href":"/blog/prompt-versioning-lifecycle-management-2026","cat":"Blog"},{"title":"Evaluating CrewAI Agents: Role Adherence Is the Unit (2026)","desc":"Evaluating CrewAI agents in 2026: role adherence as the primary metric, plus task delegation, crew coherence, and manager-worker fidelity.","href":"/blog/evaluating-crewai-agents-2026","cat":"Blog"},{"title":"Distributed Eval Runners: Celery, Ray, Temporal, Kubernetes","desc":"Celery, Ray, Temporal, and Kubernetes optimise for different things. Pick by your bottleneck, not by fashion. 2026 engineering decision guide.","href":"/blog/llm-eval-distributed-runners-2026","cat":"Blog"},{"title":"Step-by-Step Guide to MCP Evaluation (2026)","desc":"A 2026 workflow for evaluating MCP servers end to end: functional checks, security checks, cross-client compatibility, stress tests, and the CI gate.","href":"/blog/step-by-step-guide-mcp-evaluation-2026","cat":"Blog"},{"title":"What is Retrieval Augmented Generation (RAG)? The 2026 Definition","desc":"RAG is four pipeline stages and three failure modes per stage. A methodology reference for picking each stage and measuring what it broke.","href":"/blog/what-is-retrieval-augmented-generation-2026","cat":"Blog"},{"title":"Best 5 AWS Bedrock Alternatives for LLM Routing in 2026","desc":"Five AWS Bedrock alternatives for LLM routing on model catalog, cross-cloud, IAM, eval and optimizer. What each actually fixes if Bedrock is your gateway.","href":"/blog/best-aws-bedrock-llm-routing-alternatives-2026","cat":"Blog"},{"title":"Best 5 LlamaIndex Alternatives in 2026","desc":"Five LlamaIndex alternatives on retrieval portability, abstraction weight, polyglot support. What each actually fixes when RAG-first heritage stops paying.","href":"/blog/best-llamaindex-alternatives-2026","cat":"Blog"},{"title":"Evaluating LLM Summarization: A Step-by-Step Guide (2026)","desc":"Summarization eval is four rubrics, not one number: groundedness, completeness, factuality, conciseness, calibrated against humans in CI.","href":"/blog/evaluate-llm-summarization-step-by-step-2026","cat":"Blog"},{"title":"IVR Modernization: Migrate Legacy IVR to AI Voice Agents in 2026","desc":"A step-by-step IVR modernization playbook for 2026: audit legacy flows, pick a runtime, simulate, deploy, observe. Migrate DTMF menus to AI voice agents.","href":"/blog/ivr-modernization-ai-voice-agents-2026","cat":"Blog"},{"title":"Best 5 AI Guardrails for CX AI Applications in 2026","desc":"Five AI guardrails for customer support, chatbots, voice IVR, outbound, agent-assist, KB RAG. TCPA, FCC AI-voice, Moffatt, Lingo, FTC Op AI Comply.","href":"/blog/best-cx-ai-guardrails-2026","cat":"Blog"},{"title":"Evaluating Claude Code Tool Use in 2026","desc":"Evaluating Claude Code tool use in 2026: per-tool selection F1, argument fidelity, irreversibility awareness, recovery on error, on traceAI traces.","href":"/blog/evaluating-claude-code-tool-use-2026","cat":"Blog"},{"title":"A Gentle Introduction to LLM Evaluation (2026)","desc":"Learn LLM evaluation from the inside out: the three primitives (deterministic, embedding, judge), offline vs online, the starter workflow for production.","href":"/blog/gentle-introduction-llm-evaluation-2026","cat":"Blog"},{"title":"LiteLLM Compromised 2026: Incident Response and Gateway Migration","desc":"Full breakdown of the March 24 2026 LiteLLM supply chain attack: timeline, three-stage payload, detection commands, and a managed-gateway migration path.","href":"/blog/litellm-compromised-incident-response-migration-guide","cat":"Blog"},{"title":"Best Cost-Efficient AI Evaluation Platforms 2026: 7","desc":"Cost-efficient AI evaluation in 2026 is the cascade: classifiers, local heuristics, cheap judges. 7 platforms compared on per-eval cost.","href":"/blog/best-cost-efficient-ai-evaluation-platforms-2026","cat":"Blog"},{"title":"LLM-Judge Bias Mitigation (2026): Detect, Measure, Fix","desc":"Five named LLM-judge biases, each with a measurement and a mitigation that holds in production: position, verbosity, self-preference, format.","href":"/blog/evaluating-llm-judge-bias-mitigation-2026","cat":"Blog"},{"title":"How to Reduce MCP Token Costs for Claude Code at Scale in 2026","desc":"Practical 2026 how-to for cutting MCP token spend on Claude Code at fleet scale: five levers, the mcp.json + gateway config, metrics that prove the cut.","href":"/blog/how-to-reduce-mcp-token-costs-claude-code-scale-2026","cat":"Blog"},{"title":"LLM Eval vs Product Analytics: Two Layers, One Loop (2026)","desc":"Product analytics measures user behavior. LLM eval measures system behavior. 2026 PM and ML guide to keeping them separate and joining on identifier.","href":"/blog/llm-eval-vs-product-analytics-2026","cat":"Blog"},{"title":"Best AI Gateway for Replit Agent Production Workflows 2026","desc":"Five AI gateways scored on Replit Agent in 2026: per-app budgets, secret scanning, deploy-snapshot audit, multi-tenant cost slicing, tool-call survival.","href":"/blog/best-ai-gateway-replit-agent-production-workflows-2026","cat":"Blog"},{"title":"Best AI Gateway for Sweep AI and Automated Code Review Workflows in 2026","desc":"Five AI gateways scored on Sweep AI code-review workloads in 2026: per-PR cost attribution, latency budgets, false-positive loops, per-repo policy, audit.","href":"/blog/best-ai-gateway-sweep-ai-code-review-workflows-2026","cat":"Blog"},{"title":"LLM Hallucination: A 2026 Architectural Deep Dive","desc":"Hallucination is four distinct failure modes: factual, grounding, citation, reasoning. Each needs a different detector and a different fix, with code.","href":"/blog/llm-hallucination-deep-dive-2026","cat":"Blog"},{"title":"What is Evals Engineering? The Discipline Behind Production LLMs in 2026","desc":"Evals engineering is DevOps for LLMs: building, maintaining, and gating eval suites that catch real production failures. Role, tooling, and 2026 patterns.","href":"/blog/what-is-evals-engineering-2026","cat":"Blog"},{"title":"AI Agent Compliance and Governance (2026): A Runtime Playbook","desc":"Wire policy, enforcement, and audit into runtime so EU AI Act, NIST AI RMF, and ISO 42001 close on one plane without slowing releases.","href":"/blog/ai-agent-compliance-governance-2026","cat":"Blog"},{"title":"Best Claude Code Gateway for Enterprises in 2026","desc":"","href":"/blog/best-claude-code-gateway-for-enterprises-2026","cat":"Blog"},{"title":"Evaluating Vertex AI Agent Engine in 2026","desc":"Vertex ships a managed runtime. Score Vertex Search retrieval, grounded-vs-reasoning outputs, and Gemini safety filter precision/recall.","href":"/blog/evaluating-vertex-ai-agent-engine-2026","cat":"Blog"},{"title":"LLM Eval vs LLM Observability in 2026: The Disambiguation Guide","desc":"What LLM observability captures, what LLM evaluation scores, where the two overlap, and the seven axes that separate them in 2026 across vendors.","href":"/blog/llm-eval-vs-llm-observability-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Education in 2026: FERPA-Compliant Student-Data Routing","desc":"Five AI gateways for K-12, universities, and EdTech in 2026, scored on FERPA student-data routing, COPPA consent, and state privacy laws.","href":"/blog/best-ai-gateways-education-2026","cat":"Blog"},{"title":"Best 5 Aporia Alternatives in 2026","desc":"Five Aporia alternatives on inline guardrail latency, eval-loop wiring, deployment, SDK breadth. What each actually fixes after the Coralogix integration.","href":"/blog/best-aporia-alternatives-2026","cat":"Blog"},{"title":"Best 5 Voice AI Simulation Tools for Hospitality in 2026","desc":"Five voice AI simulation tools for hospitality, hotel reservation, airline rebooking, multi-lingual concierge. ADA Title III, PCI DSS, TCPA, DOT.","href":"/blog/best-hospitality-voice-ai-simulation-2026","cat":"Blog"},{"title":"LLM Agent Evaluation: The Complete Guide (2026)","desc":"Agent eval guide built on the closed loop: offline eval, CI gate, production trace eval, Error Feed, and optimization. Treat eval as a loop.","href":"/blog/llm-agent-evaluation-complete-guide-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Government in 2026: FedRAMP-Ready Gateways With Audit Trails","desc":"Five AI gateways for federal, DoD, and state government 2026, scored on FedRAMP Moderate and High, DoD IL2 to IL6, NIST AI RMF, EO 14110.","href":"/blog/best-ai-gateways-government-2026","cat":"Blog"},{"title":"Best 5 CrewAI Alternatives in 2026","desc":"Five CrewAI alternatives on framework mental model, multi-agent ergonomics, API stability. What each actually fixes when CrewAI prototypes hit production.","href":"/blog/best-crewai-alternatives-2026","cat":"Blog"},{"title":"Edge Cases and Adversarial Inputs in LLM Evaluation (2026)","desc":"Systematically generate and evaluate edge cases plus adversarial inputs for LLM agents in 2026: seven categories, five generation methods, five-step plan.","href":"/blog/llm-eval-edge-cases-adversarial-2026","cat":"Blog"},{"title":"TruLens vs Future AGI in 2026: Feedback-Function Library vs Eval-Stack Package","desc":"TruLens vs Future AGI compared honestly on feedback functions, Snowflake fit, cost, runtime guardrails, distributed runners, closed loop. Where each wins.","href":"/blog/trulens-vs-future-agi-2026","cat":"Blog"},{"title":"Best 5 H2O AI Cloud LLM Alternatives in 2026","desc":"Five H2O AI Cloud LLM alternatives on purpose-built tooling, gateway and routing depth, pricing decoupled from H2O bundles, what each actually fixes.","href":"/blog/best-h2o-ai-cloud-llm-alternatives-2026","cat":"Blog"},{"title":"Evaluating Modal LLM Inference Apps in 2026","desc":"Evaluate Modal-served LLM apps in 2026: per-type latency parity (cold, warm, concurrent), p99 tail quality under burst, shutdown determinism.","href":"/blog/evaluating-modal-llm-inference-2026","cat":"Blog"},{"title":"Real-Time STT vs Offline STT: A 2026 Decision Guide for Voice AI","desc":"Real-time STT vs offline STT in 2026: latency, WER, cost, accent robustness, and the eval rubric that scores both at scale. A decision matrix for voice AI.","href":"/blog/real-time-stt-vs-offline-stt-2026","cat":"Blog"},{"title":"Voice AI Infrastructure Stack: A 2026 Pick-by-Stage Decision Guide","desc":"Pick-by-stage guide to the 2026 voice AI stack: telephony, orchestration, STT, LLM, TTS, eval, observe. Real pricing, real picks, three full compositions.","href":"/blog/voice-ai-infrastructure-stack-decision-guide-2026","cat":"Blog"},{"title":"Evaluating Browser-Use Agents in 2026: The Six Failure Modes","desc":"Evaluating browser-use agents in 2026: WebArena grades happy-path completion; production grades recovery from six failure modes nobody benchmarks.","href":"/blog/evaluating-browser-use-agents-2026","cat":"Blog"},{"title":"Langfuse Alternatives in 2026: 5 Honest Picks for Production AI","desc":"Honest 2026 comparison of Langfuse alternatives: Future AGI, LangSmith, Phoenix, Braintrust, Helicone on eval depth, gateway, and the loop.","href":"/blog/langfuse-alternatives-2026","cat":"Blog"},{"title":"LLM Arena as a Judge: Pairwise Comparison Evals (2026)","desc":"Arena-as-a-judge in 2026: when pairwise wins over rubric scoring, the four biases to control, the CI math, and the wiring back to production OTel traces.","href":"/blog/llm-arena-judge-comparison-evals-2026","cat":"Blog"},{"title":"What Does a Good LLM Trace Look Like in 2026: Anatomy and Attributes","desc":"Anatomy of a good LLM trace in 2026: span hierarchy, OTel GenAI attributes, prompt-version tags, eval scores, cost attribution, retrieval and tool spans.","href":"/blog/what-does-a-good-llm-trace-look-like-2026","cat":"Blog"},{"title":"Best AI Gateway to Manage Codex CLI Token Spend in 2026","desc":"Five AI gateways for Codex CLI token-spend management in 2026: per-session attribution, per-dev caps, alerts, model downgrade, cache observability.","href":"/blog/best-ai-gateway-manage-codex-cli-token-spend-2026","cat":"Blog"},{"title":"Best 5 New Relic AI Monitoring Alternatives in 2026","desc":"Five New Relic AI Monitoring alternatives ranked on LLM-native depth, ingest-based pricing curves, eval support, gateway and routing, APM-tool fixes.","href":"/blog/best-new-relic-ai-monitoring-alternatives-2026","cat":"Blog"},{"title":"Best 5 OctoML Alternatives for LLM Inference in 2026","desc":"Five OctoML alternatives scored on hosted-inference depth, throughput, fine-tuning, infra control after NVIDIA narrowed OctoML to model compilation.","href":"/blog/best-octoml-llm-inference-alternatives-2026","cat":"Blog"},{"title":"LLM Eval vs RLHF Feedback Loops in 2026","desc":"How production LLM eval feeds RLHF, RLAIF, and DPO preference tuning: five feedback-loop patterns, six-step eval-driven pipeline, when post-training wins.","href":"/blog/llm-eval-vs-rlhf-feedback-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Compliance Audit Trails in 2026","desc":"Five AI gateways on what compliance officers actually need: tamper-evident logs, 3-10 yr retention, SIEM exports, span granularity, legal hold, frameworks.","href":"/blog/best-ai-gateways-compliance-audit-trails-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Opencode Token Tracking and Access Controls in 2026","desc":"Five AI gateways scored on Opencode token tracking and access controls in 2026: per-dev attribution, per-repo budgets, audit logs, self-host posture, gaps.","href":"/blog/best-ai-gateways-opencode-token-tracking-2026","cat":"Blog"},{"title":"Frontier Model Safety Analysis (2026): RSP, Preparedness, FSF","desc":"Anthropic RSP, OpenAI Preparedness, DeepMind FSF: where frontier-lab safety frameworks converge, diverge, and how enterprises operationalize them.","href":"/blog/frontier-model-safety-analysis-2026","cat":"Blog"},{"title":"The LLM Eval Vendor Buyer's Guide for 2026","desc":"Heads-of-engineering buyer guide for LLM eval vendors 2026. Ten criteria, eight vendor categories scored honestly, 5-question rubric, procurement flow.","href":"/blog/llm-eval-vendor-buyer-guide-2026","cat":"Blog"},{"title":"Evaluating Anyscale Ray Serve LLM Apps in 2026","desc":"How to evaluate an Anyscale Ray Serve LLM in 2026: catch autoscaling lag, replica skew, and tail-quality cliffs the model eval never sees.","href":"/blog/evaluating-anyscale-ray-llm-2026","cat":"Blog"},{"title":"Evaluating Embedding Models in 2026","desc":"MTEB Recall@10 does not transfer to your domain. 500 labeled query-passage pairs from your traffic decide which embedding wins, not boards.","href":"/blog/evaluating-embedding-models-2026","cat":"Blog"},{"title":"LLM Eval vs Software Testing: The 2026 Bridge for Dev Teams","desc":"How engineers should map the test pyramid (unit, integration, e2e) onto LLM eval in 2026: the seven gaps, the analogy, a five-step transition.","href":"/blog/llm-eval-vs-software-testing-2026","cat":"Blog"},{"title":"Multi-Turn Jailbreaking (Defender's Guide 2026)","desc":"Single-turn guardrails lose to multi-turn adversaries. Crescendo, Cipher, role lock-in, many-shot ICL succeed. The defense stack that catches.","href":"/blog/multi-turn-jailbreaking-defender-2026","cat":"Blog"},{"title":"Best AI Gateway for Continue.dev VSCode Workflow in 2026","desc":"Five AI gateways scored on the Continue.dev VSCode workflow in 2026: autocomplete latency, chat sessions, per-user attribution, hybrid routing.","href":"/blog/best-ai-gateway-continue-dev-vscode-workflow-2026","cat":"Blog"},{"title":"Evaluating LangChain RAG Applications in 2026","desc":"LangChain RAG eval is two problems: the retriever and the chain. Per-step rubrics catch the bug; chain-level Groundedness on LCEL output confirms the fix.","href":"/blog/evaluating-langchain-rag-2026","cat":"Blog"},{"title":"The State of LLM Benchmarking (2026): What the Leaderboards Tell You, and What They Don't","desc":"MMLU, GSM8K, SWE-bench Verified, BFCL, tau-bench, GPQA, ARC-AGI-2, Chatbot Arena. What each measures, where each breaks, triangulate-plus-private 2026.","href":"/blog/llm-benchmarking-state-2026","cat":"Blog"},{"title":"Prompt Regression Testing: A Practical 2026 Guide","desc":"Prompt regression is pytest for prompts. Three patterns: per-rubric assertion, per-route stratified eval, paired comparison vs prior version with CI delta.","href":"/blog/prompt-regression-testing-2026","cat":"Blog"},{"title":"How I Built a Deterministic LLM Evaluation Library","desc":"A first-person write-up: why a $40K judge bill pushed me to build deterministic LLM evaluation metrics first, schema, regex, structural, citation-validity.","href":"/blog/how-i-built-deterministic-llm-evaluation-metrics-2026","cat":"Blog"},{"title":"Academic vs Production LLM Evaluation: The 2026 Bridge","desc":"Academic LLM benchmarks answer which model is smartest. Production eval answers if your system works on your traffic. Methodologies and bridge patterns.","href":"/blog/llm-eval-academic-vs-production-2026","cat":"Blog"},{"title":"How to Organize an LLM Eval Team in 2026: Three-Role Split, RACI, Hiring","desc":"Eval ownership splits three ways: platform team owns tooling, product teams own rubrics, quality council owns policy. RACI, org models, hiring, FAGI.","href":"/blog/llm-eval-team-organization-2026","cat":"Blog"},{"title":"What Is a Fallback Strategy for LLM APIs in 2026?","desc":"A 2026 field guide to LLM fallback strategy: definition, five strategies, architecture, buyer's guide, myths, and OTel GenAI auditable spans.","href":"/blog/what-is-llm-fallback-strategy-2026","cat":"Blog"},{"title":"5 Best AI Appointment Booking Voice Tools in 2026","desc":"Five AI appointment booking voice tools ranked for 2026: Vapi, Retell, Synthflow, Bland, Goodcall on calendars, latency, reliability.","href":"/blog/best-ai-appointment-booking-voice-tools-2026","cat":"Blog"},{"title":"Best 5 AI Gateways to Run Claude Code with Any LLM Provider in 2026","desc":"Five AI gateways scored on running Claude Code against non-Anthropic models in 2026: translation fidelity, tool-use survival, streaming, latency, routing.","href":"/blog/best-ai-gateways-claude-code-any-llm-provider-2026","cat":"Blog"},{"title":"Best 5 Guardrails AI Alternatives in 2026","desc":"Five Guardrails AI alternatives on inline runtime latency, native gateway, eval, multi-language. What each actually fixes beyond Python-only validation.","href":"/blog/best-guardrails-ai-alternatives-2026","cat":"Blog"},{"title":"7 Best TTS Providers for Voice AI Agents in 2026 (Tested + Ranked)","desc":"Tested and ranked: 7 best TTS providers for voice AI agents 2026, with real per-character pricing, streaming TTFA latency, voice cloning, SSML support.","href":"/blog/best-tts-providers-voice-agents-2026","cat":"Blog"},{"title":"Future AGI vs Hamming: 2026 Voice Agent Testing Comparison","desc":"Future AGI vs Hamming on eval rubrics, native voice observability, simulation, guardrails, optimization, compliance. Where each actually fits in 2026.","href":"/blog/future-agi-vs-hamming-2026","cat":"Blog"},{"title":"HIPAA-Compliant Voice AI in 2026: Build, Test, Deploy","desc":"End-to-end HIPAA voice AI in 2026. BAA-covered call chain, PHI-aware regression suite, breach detection, patient-access flows, with Future AGI Protect.","href":"/blog/hipaa-compliant-voice-ai-build-test-deploy-2026","cat":"Blog"},{"title":"How to Evaluate TTS Quality for Voice AI in 2026: SSML + MOS + Rubrics","desc":"Evaluate TTS quality for voice AI in 2026 with audio_quality rubrics, MOS scoring, SSML snapshot regression, and A/B provider comparison via Future AGI.","href":"/blog/how-to-evaluate-tts-quality-voice-ai-2026","cat":"Blog"},{"title":"How to Monitor AI Voice Agents in Production: 2026 Playbook","desc":"Two-category playbook for monitoring AI voice agents: native FAGI dashboard for Vapi-class, traceAI SDK for Pipecat and LiveKit, plus SLOs.","href":"/blog/how-to-monitor-ai-voice-agents-production-2026","cat":"Blog"},{"title":"Multilingual Voice AI Testing: A 2026 Engineering Guide","desc":"Engineer multilingual voice AI tests across many languages. Translation_accuracy, cultural_sensitivity, custom evaluators, ElevenLabs voice coverage.","href":"/blog/multilingual-voice-ai-testing-2026","cat":"Blog"},{"title":"An Introduction to Production Monitoring for Voice Agents in 2026","desc":"What production monitoring means for voice agents in 2026: definitions, what changes vs text, a reference architecture, and a getting-started checklist.","href":"/blog/production-monitoring-voice-agents-intro-2026","cat":"Blog"},{"title":"Sub-500ms Voice AI: The Complete Latency Budget Guide for 2026","desc":"How to hit a sub-500ms P95 voice AI turn in 2026. Per-stage budget, engineering choices, when sub-500ms is the right target and when it is not.","href":"/blog/sub-500ms-voice-ai-guide-2026","cat":"Blog"},{"title":"Voice AI for Banking and Financial Services in 2026","desc":"How to deploy voice AI across banking workflows in 2026. Account servicing, fraud verification, loan qualification, payment processing, dispute resolution.","href":"/blog/voice-ai-banking-financial-services-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Aider with Local Models in 2026","desc":"Five AI gateways scored on Aider with local models in 2026: OpenAI-compatible passthrough to Ollama and vLLM, hosted fallback, GPU-aware routing, gaps.","href":"/blog/best-ai-gateways-aider-local-models-2026","cat":"Blog"},{"title":"Evaluate Google ADK Agents: 6-Step 2026 Production Loop","desc":"Evaluate Google ADK agents in 6 steps: traceAI instrumentation, span-attached evaluate() scoring, AgentEvaluator CI gates, persona sim, Bayesian opt.","href":"/blog/evaluate-google-adk-agents","cat":"Blog"},{"title":"Evaluating LLM Context Window Management (2026)","desc":"Long-context support is marketing. Long-context fidelity is what you eval: NIAH at every position, lost-in-middle on your docs, attention-budget cost.","href":"/blog/evaluating-llm-context-window-management-2026","cat":"Blog"},{"title":"Evaluating Multi-Turn Conversations: A Deep Dive (2026)","desc":"Multi-turn eval is not turn-by-turn averaged. The unit is the trajectory: coherence, context, intent drift, resolution. The playbook.","href":"/blog/evaluating-multi-turn-conversations-deep-dive-2026","cat":"Blog"},{"title":"What Is LLM Routing? A 2026 Field Guide","desc":"A 2026 field guide to LLM routing: definition, five strategies (round-robin, weighted, latency, cost, quality), architecture, myths.","href":"/blog/what-is-llm-routing-2026","cat":"Blog"},{"title":"Best 5 Deepchecks Alternatives in 2026","desc":"Five Deepchecks alternatives scored on LLM-native evaluators, language coverage, pricing, gateway depth, community, for agent workloads.","href":"/blog/best-deepchecks-alternatives-2026","cat":"Blog"},{"title":"Best 5 OpenLIT Alternatives in 2026","desc":"Five OpenLIT alternatives on community, evaluator depth, hosted-dashboard maturity, gateway, optimizer. What each actually fixes beyond OTel-only OSS.","href":"/blog/best-openlit-alternatives-2026","cat":"Blog"},{"title":"Evaluating Haystack RAG Pipelines in 2026","desc":"Haystack Pipelines are component DAGs, not black boxes. Per-component rubrics on Retriever, Ranker, Generator + pipeline-level Groundedness.","href":"/blog/evaluating-haystack-rag-2026","cat":"Blog"},{"title":"Evaluating LLM PII Detection (2026)","desc":"PII detection eval is per-entity precision AND recall on adversarial AND benign sets. One F1 score hides a HIPAA breach. The 2026 methodology.","href":"/blog/evaluating-llm-pii-detection-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for E-commerce in 2026: Search, Personalization, and Checkout","desc":"Five AI gateways for pure-play e-commerce in 2026: product-search recall, recommendation lift, cart latency, GDPR/CCPA consent, PCI-DSS, EU DSA scope.","href":"/blog/best-ai-gateways-ecommerce-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Media and Publishing in 2026: Content Routing, Licensing, and Brand-Safety Guardrails","desc":"Five AI gateways for media and publishing 2026: source-attribution, copyright scoring, brand safety, deepfake detection, C2PA watermarking.","href":"/blog/best-ai-gateways-media-publishing-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Prompt Management in 2026","desc":"Five AI gateways for prompt management in 2026 scored on version pinning, per-template A/B split, sub-60s rollback, variable safety, eval-gated promotion.","href":"/blog/best-ai-gateways-prompt-management-2026","cat":"Blog"},{"title":"The LLM Evaluation Glossary (2026 Definitions)","desc":"A practitioner's dictionary for LLM evaluation in 2026: the 30 most-confused terms, what each means, and the adjacent terms it conflates with.","href":"/blog/llm-evaluation-glossary-definitions-2026","cat":"Blog"},{"title":"Best 5 Flowise Alternatives in 2026","desc":"Five Flowise alternatives on canvas ergonomics, scale beyond the visual builder, ecosystem. What each actually fixes when drag-and-drop stops working.","href":"/blog/best-flowise-alternatives-2026","cat":"Blog"},{"title":"Best 5 HumanSignal Alternatives in 2026","desc":"Five HumanSignal/Label Studio alternatives on labeled-data portability, runtime LLM eval, guardrails, optimizer. What actually fixes production agents.","href":"/blog/best-humansignal-alternatives-2026","cat":"Blog"},{"title":"Evaluating LLM Structured Output Modes (2026)","desc":"Compare OpenAI strict, Anthropic JSON, Gemini schema, and Outlines grammar-constrained generation: schema-validity rate, quality tax, failure modes.","href":"/blog/evaluating-llm-structured-output-modes-2026","cat":"Blog"},{"title":"Evaluating Prompt Caching Quality in 2026","desc":"Prompt caching saves 50-90% on spend but ships two silent regressions: invalidation bugs and semantic-cache wrong-prompt hits. The eval that catches both.","href":"/blog/evaluating-prompt-caching-quality-2026","cat":"Blog"},{"title":"Evaluating AutoGen Agents: The Handoff Is the Unit (2026)","desc":"Evaluating AutoGen agents in 2026: the handoff is the eval unit. Three failure modes, three rubrics, per-pair spans, and the production loop.","href":"/blog/evaluating-autogen-agents-2026","cat":"Blog"},{"title":"LLM Eval Myths: Six Skeptical Objections, Honestly Answered (2026)","desc":"Six skeptical objections to LLM eval. Five are right about something the field undersells, one is laziness. Honest answers to each, in turn.","href":"/blog/llm-eval-myths-skeptics-2026","cat":"Blog"},{"title":"LLM Evaluation in 2027-2028: Ten Predictions for Eval-Stack Buyers","desc":"Where LLM evaluation is heading in 2027 and 2028: ten grounded predictions on cost, classifiers, self-improving judges, CI gates, FinOps, multimodal.","href":"/blog/llm-evaluation-future-predictions-2027-2028","cat":"Blog"},{"title":"Top 5 G-Eval Use Cases (2026): Where Rubric Judges Actually Win","desc":"Five use cases where G-Eval is the right primitive: subjective rubric scoring, faithfulness on free-form text, custom-domain rubrics, weighted, reasoning.","href":"/blog/top-5-geval-use-cases-2026","cat":"Blog"},{"title":"Best 5 RAG Evaluation Tools for HR AI Applications in 2026","desc":"Five RAG eval tools for HR, benefits Q&A, policy lookup, manager-toolkit, leave-eligibility. NYC AEDT, EEOC, Mobley v Workday, EU AI Act III, AB 2930.","href":"/blog/best-hr-rag-evaluation-2026","cat":"Blog"},{"title":"Evaluating Coding Agents 2026: A Five-Dimension Eval","desc":"Public SWE-bench scores don't transfer to your codebase. Evaluate coding agents on golden-PR replay, tool calls, multi-file coherence, plan, and rollback.","href":"/blog/evaluating-coding-agents-2026","cat":"Blog"},{"title":"How to Optimize LiveKit Voice Agent Latency in 2026: 12 Techniques + Code","desc":"Cut LiveKit Agents voice latency to sub-500ms p95 in 2026. 12 techniques with real AgentSession code: streaming STT, partial TTS, prefix caching, regional.","href":"/blog/how-to-optimize-livekit-latency-2026","cat":"Blog"},{"title":"RAG Evaluation Metrics: A Deep Dive (2026)","desc":"RAG eval is a bisection problem. Each metric pins failure to retrieval, generation, or the cross-cut. Metric-by-metric methodology guide.","href":"/blog/rag-evaluation-metrics-deep-dive-2026","cat":"Blog"},{"title":"Best 5 Halluminate Alternatives in 2026","desc":"Five Halluminate alternatives on evaluator breadth, runtime guardrails, self-host, languages. What each actually fixes beyond hallucination-only detection.","href":"/blog/best-halluminate-alternatives-2026","cat":"Blog"},{"title":"How to Connect Claude Code to an MCP Gateway in 2026","desc":"Wiring Claude Code to an MCP gateway 2026: mcp.json config, routing rules, per-server auth scoping, verification. Production checklist and gateway picks.","href":"/blog/how-to-connect-claude-code-to-mcp-gateway-2026","cat":"Blog"},{"title":"LLM Safety and AI Regulations (2026): What's Binding, What's Not","desc":"EU AI Act, NIST AI RMF, ISO 42001, India DPDPA, US executive orders: which AI safety rules are binding mid-2026, and how to wire eval to them.","href":"/blog/llm-safety-ai-regulations-2026","cat":"Blog"},{"title":"Voice AI Observability for Retell AI: 2026 Implementation Guide","desc":"Wire Retell AI observability the FAGI way: native dashboard via Assistant ID, optional traceAI SDK, eval engine on every call with audio + transcript.","href":"/blog/voice-ai-observability-retell-2026","cat":"Blog"},{"title":"Best AI Agent Guardrails Platforms in 2026: 6 Tools Compared","desc":"Senior-engineer comparison of the best AI agent guardrails platforms 2026: latency budgets, placement, deployment topology, vendor fit.","href":"/blog/best-ai-agent-guardrails-platforms-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Prompt Injection Defense in 2026","desc":"Five AI gateways for prompt injection defense in 2026 scored on direct and indirect detection, sub-100ms inline latency, openness, modes, robustness.","href":"/blog/best-ai-gateways-prompt-injection-defense-2026","cat":"Blog"},{"title":"LLM Eval for Enterprises in 2026: The F500 Playbook","desc":"The enterprise LLM evaluation playbook for Fortune 500 rollouts: multi-BU governance, regulatory rubric mapping, data residency, chargeback, procurement.","href":"/blog/llm-eval-for-enterprises-2026","cat":"Blog"},{"title":"What is Tree of Thoughts Prompting? Branching Reasoning in 2026","desc":"Tree of Thoughts prompts an LLM to explore reasoning branches under an evaluator and search policy. When it pays off vs CoT, 2026 production patterns.","href":"/blog/what-is-tree-of-thoughts-prompting-2026","cat":"Blog"},{"title":"Evaluating LLM Self-Reflection Loops: The 3 Metrics That Matter (2026)","desc":"Self-reflection loops sometimes improve outputs and sometimes destroy them. The pre-vs-post delta, over-correction rate, and cost-per-improvement metrics.","href":"/blog/evaluating-llm-self-reflection-loops-2026","cat":"Blog"},{"title":"LLM Eval Data Drift Detection: Three Drifts That Age Your Golden Set","desc":"Eval dataset drift is the silent killer. A 2026 method for catching input, prompt-template, and retrieval-corpus drift before CI is wrong.","href":"/blog/llm-eval-data-drift-detection-2026","cat":"Blog"},{"title":"LLM Spend and Cost Tracking: Cost-per-Outcome in 2026","desc":"Cost-per-token is theater. The metric is cost-per-outcome. Per-trace attribution, gateway budget enforcement, the 2026 LLM FinOps playbook.","href":"/blog/llm-spend-cost-tracking-2026","cat":"Blog"},{"title":"OWASP LLM Top 10 (2025): Risks, Mitigations, and the Tools","desc":"OWASP LLM Top 10 (2025) for engineers: each risk, threat model, concrete mitigations, and the eval and guardrail tools that actually implement them.","href":"/blog/owasp-llm-top-10-2025-risks-mitigations-2026","cat":"Blog"},{"title":"Best 5 AutoGen Alternatives in 2026","desc":"Five AutoGen alternatives on production fit, API stability, gateway, observability, governance. What each actually fixes when MS Research stops paying.","href":"/blog/best-autogen-alternatives-2026","cat":"Blog"},{"title":"Best 5 Voice AI Simulation Tools for Insurance in 2026","desc":"Five voice AI simulation tools for insurance, FNOL intake, claims-status, fraud verification, policy Q&A. NAIC, NY Reg 187, state DOI exam authority.","href":"/blog/best-insurance-voice-ai-simulation-2026","cat":"Blog"},{"title":"Best 5 LiteLLM Alternatives in 2026","desc":"Five LiteLLM alternatives on supply-chain posture, self-host burden, UI polish, optimization. What each actually fixes after the March 2026 PyPI incident.","href":"/blog/best-litellm-alternatives-2026","cat":"Blog"},{"title":"The Eval ROI Business Case: Show the Spreadsheet, Not the Slide Deck","desc":"Eval ROI is four terms: avoided-incident, faster-ship, model-substitution, minus infra. Teams underestimate incidents 10x. With math.","href":"/blog/llm-eval-roi-business-case-2026","cat":"Blog"},{"title":"Evaluating LLM Batch Inference in 2026","desc":"Batch APIs cut LLM cost ~50%, but break the eval loop. The working pattern for deferred execution, batch-vs-sync drift, and failed-row recovery.","href":"/blog/evaluating-llm-batch-inference-2026","cat":"Blog"},{"title":"Evaluating LLM Routing Policies in 2026","desc":"Routing-policy eval is not model eval. The 2026 playbook: route correctness, realized vs theoretical cost, quality under substitution, fallback checks.","href":"/blog/evaluating-llm-routing-policies-2026","cat":"Blog"},{"title":"How to Optimize Retell Voice Agent Latency in 2026: 12 Techniques + Code","desc":"Optimize Retell AI voice agent latency to sub-500ms p95 in 2026. 12 techniques with real Retell config: STT, response_engine, backchannel, async eval.","href":"/blog/how-to-optimize-retell-latency-2026","cat":"Blog"},{"title":"Simulated Multi-Turn Conversation Eval (2026)","desc":"Build a simulated multi-turn eval that catches real failures: the Persona-Scenario-Adversary triangle, FAGI simulate-sdk patterns, trajectory scoring.","href":"/blog/simulated-multi-turn-conversation-eval-2026","cat":"Blog"},{"title":"Evaluating Fireworks AI and Together AI Inference in 2026","desc":"Evaluate Fireworks AI, Together AI, Modal, and Replicate apps in 2026: bit-fidelity, per-provider arena-judge, latency parity, quantization.","href":"/blog/evaluating-fireworks-together-inference-2026","cat":"Blog"},{"title":"Evaluating LLM Systems: Metrics and Benchmarks (2026)","desc":"Benchmarks tell you which model is smartest. Metrics tell you if your system works. 2026 guide: benchmark map, metric catalog, CI gate, rubric.","href":"/blog/evaluating-llm-systems-metrics-benchmarks-2026","cat":"Blog"},{"title":"Open Source LLM Red Team Frameworks Compared (2026)","desc":"OSS LLM red-team splits three ways: orchestrators (PyRIT), probe libraries (garak), benchmark suites (HarmBench, JailbreakBench, AdvBench).","href":"/blog/open-source-llm-red-team-frameworks-compared-2026","cat":"Blog"},{"title":"BLEU vs ROUGE vs BERTScore: Worked Examples and 2026 Use Cases","desc":"BLEU, ROUGE, BERTScore decoded with worked examples. What each measures, when each breaks, and where LLM-judge scoring replaces them in 2026.","href":"/blog/what-is-bleu-rouge-bertscore-2026","cat":"Blog"},{"title":"Best 5 Datadog LLM Observability Alternatives in 2026","desc":"Five Datadog LLM Observability alternatives on OpenInference, bundle-free pricing, gateway-native routing. What each actually fixes when you re-point OTel.","href":"/blog/best-datadog-llm-observability-alternatives-2026","cat":"Blog"},{"title":"Best 5 Enkrypt AI Alternatives in 2026","desc":"Five Enkrypt AI alternatives on inline-guardrail latency, gateway, routing, language SDK breadth. What each actually fixes when red-teaming is not enough.","href":"/blog/best-enkrypt-ai-alternatives-2026","cat":"Blog"},{"title":"Evaluating LLM Tool Use in 2026: The Four-Step Contract","desc":"Evaluating LLM tool use is a four-step contract: decide to call, pick the tool, build args, integrate the result. Score each step independently.","href":"/blog/evaluating-llm-tool-use-2026","cat":"Blog"},{"title":"LLM App Observability with OpenTelemetry: The 2026 Setup","desc":"OTel for LLM apps in 2026 = OTel-GenAI + OpenInference + eval-as-span-attribute. Three layers, traceAI register pattern, span enrichment, sampling.","href":"/blog/llm-app-observability-otel-2026","cat":"Blog"},{"title":"Top Enterprise AI Gateways for Governing Claude Code in 2026","desc":"Five enterprise AI gateways scored on governing Claude Code in 2026: model whitelists, repo allowlists, approval workflows, policy injection, audit trails.","href":"/blog/top-enterprise-ai-gateways-governing-claude-code-2026","cat":"Blog"},{"title":"traceAI: OpenTelemetry LLM Tracing in 2 Lines of Code","desc":"Open-source Apache 2.0 OpenTelemetry tracing for LLM apps: 50+ AI surfaces across Python, TypeScript, Java, C#. Two lines, zero lock-in.","href":"/blog/traceai-opentelemetry-llm-tracing","cat":"Blog"},{"title":"AI Conversation Monitoring for Voice Agents: 6 Metrics in 2026","desc":"Monitor voice agent conversations with 6 metrics in 2026: turn coherence, intent confidence, completion, sentiment, escalation, and repeat-question signal.","href":"/blog/voice-agent-conversation-monitoring-2026","cat":"Blog"},{"title":"What Is Semantic Caching for LLMs? A 2026 Guide","desc":"Canonical 2026 semantic caching for LLMs: exact vs semantic, embeddings, threshold, TTL, invalidation, 5 patterns production teams actually ship.","href":"/blog/what-is-semantic-caching-llms-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Healthcare in 2026: HIPAA-Ready Routing With Built-In PHI Guardrails","desc":"Five AI gateways for healthcare in 2026, scored on HIPAA BAA coverage, PHI redaction, and the 2026 compliance stack: HIPAA NPRM, HTI-1 DSI, FDA PCCP.","href":"/blog/best-ai-gateways-healthcare","cat":"Blog"},{"title":"Best 5 KServe Alternatives for LLM Inference in 2026","desc":"Five KServe alternatives for LLM workloads, scored on serving throughput, K8s-native fit, gateway and observability surfaces, and migration plan.","href":"/blog/best-kserve-llm-alternatives-2026","cat":"Blog"},{"title":"Best 5 Mistral La Plateforme Alternatives in 2026","desc":"Five Mistral La Plateforme alternatives on multi-provider routing, gateway depth, eval, optimizer. What each actually fixes when you outgrow Mistral-only.","href":"/blog/best-mistral-la-plateforme-alternatives-2026","cat":"Blog"},{"title":"LLM Deployment Best Practices in 2026: A Production Checklist","desc":"LLM deployment in 2026: traceAI, OTel, prompt versioning, eval gates, guardrails, gateway routing, fallback patterns. The production checklist that ships.","href":"/blog/llm-deployment-best-practices-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Legal in 2026: Privilege, Citation Tracking, and SOC 2","desc":"Five AI gateways for AmLaw 100 and corporate legal in 2026 scored on privilege isolation, citation verification, SOC 2 Type II/III, on-prem, matter cost.","href":"/blog/best-ai-gateways-legal-2026","cat":"Blog"},{"title":"Deterministic vs LLM-Judge Evals (2026): Layer, Don't Choose","desc":"Deterministic vs LLM-judge isn't a pick. It's a cascade. Where each wins, where each breaks, and the layering that drops eval cost 95% in production.","href":"/blog/deterministic-vs-llm-judge-evals-2026","cat":"Blog"},{"title":"Evaluating LLM Confidence and Uncertainty (2026)","desc":"Logprob aggregation, semantic entropy, Brier score, and Platt scaling. The 2026 methodology for calibrated LLM confidence scores you can actually trust.","href":"/blog/evaluating-llm-confidence-uncertainty-2026","cat":"Blog"},{"title":"LLM Model Bias and Fairness Evaluation (2026)","desc":"Evaluating output-side bias in LLMs: seven fairness axes, four measurement techniques, regulatory frame, and FAGI surfaces that make audits continuous.","href":"/blog/llm-eval-bias-fairness-2026","cat":"Blog"},{"title":"Best 5 MCP Gateways in 2026: Post-RCE Production Picks","desc":"Five MCP gateways for production AI agents in 2026, scored on the Future AGI Production Gateway Scorecard after the April Anthropic STDIO RCE.","href":"/blog/best-mcp-gateways","cat":"Blog"},{"title":"Evaluating Agentic Workflows Orchestrated With Temporal","desc":"Temporal turns agent workflows into replayable state machines. Eval per activity, workflow outcome, retry budget, signal-handler correctness.","href":"/blog/evaluating-agentic-workflows-temporal-2026","cat":"Blog"},{"title":"Evaluating Cheap Frontier Models in 2026","desc":"How to evaluate DeepSeek-V3, Qwen, Llama 3.3, Mistral, and Phi-4 as production substitutes for Claude and GPT-5 without a silent quality cliff.","href":"/blog/evaluating-cheap-frontier-models-2026","cat":"Blog"},{"title":"traceAI: OpenTelemetry-Native LLM and Agent Tracing in 2026","desc":"traceAI is the open-source OpenTelemetry-native tracing library for LLM and agent apps. Span model, 30+ integrations, OTLP transport, how to choose.","href":"/blog/traceai-opentelemetry-tracing-2026","cat":"Blog"},{"title":"Evaluating Deep Research Agents in 2026","desc":"Aggregate quality hides which research stage broke. Score plan, retrieve, source, claim, and synthesis independently or you cannot fix anything.","href":"/blog/evaluating-deep-research-agents-2026","cat":"Blog"},{"title":"Evaluating Google ADK Agents in 2026","desc":"Google ADK's opinionated primitives (Sequential, Parallel, Loop, sub-agent dispatch) demand ADK-native eval, not a LangChain rig in a trench coat.","href":"/blog/evaluating-google-adk-agents-2026","cat":"Blog"},{"title":"Evaluating OpenAI Agents SDK: The Handoff Is the Test (2026)","desc":"Evaluating OpenAI Agents SDK in 2026: handoff correctness, output_type schema fidelity, guardrail invocation, tool-call accuracy across four primitives.","href":"/blog/evaluating-openai-agents-sdk-2026","cat":"Blog"},{"title":"Evaluating RAG Faithfulness: A 2026 Deep Dive","desc":"Why answer-level Groundedness hides RAG hallucinations, and how claim-level decomposition, cherry-pick detection, and sycophancy scoring fix it.","href":"/blog/evaluating-rag-faithfulness-deep-dive-2026","cat":"Blog"},{"title":"Best AI Gateway for Tabnine Enterprise in 2026","desc":"","href":"/blog/best-ai-gateway-tabnine-enterprise-2026","cat":"Blog"},{"title":"Best 5 AI Gateways to Govern GitHub Copilot in the Enterprise in 2026","desc":"Five AI gateways scored on enterprise GitHub Copilot governance in 2026: SSO-enforced attribution, DLP on code egress, per-repo budgets, SOX/SOC 2 audit.","href":"/blog/best-ai-gateways-github-copilot-enterprise-2026","cat":"Blog"},{"title":"LLM Eval Stack: A Reference Architecture for 2026","desc":"The 8-layer LLM eval reference architecture for 2026: ASCII diagrams end to end, five deployment topologies, integration points, anti-patterns it kills.","href":"/blog/llm-eval-stack-reference-architecture-2026","cat":"Blog"},{"title":"What Is LLM Observability? The Ultimate 2026 Guide","desc":"LLM observability in 2026 is OpenTelemetry plus LLM-aware spans plus eval-as-span-attribute. The reference guide for ML engineers picking a stack.","href":"/blog/what-is-llm-observability-ultimate-guide-2026","cat":"Blog"},{"title":"Future AGI vs TrueFoundry in 2026: AI-Native Loop vs MLOps Bundle","desc":"Future AGI vs TrueFoundry on routing, observability, cost attribution, security, deployment, DX. Honest verdict, pricing, where each falls short in 2026.","href":"/blog/future-agi-vs-truefoundry-2026","cat":"Blog"},{"title":"5 Best AI Voice Agent Platforms for Outbound Sales in 2026","desc":"5 AI voice agent platforms for outbound sales in 2026. Vapi, Retell, Synthflow, Bland, Goodcall scored on dialer depth, latency, reliability.","href":"/blog/best-ai-voice-agent-platforms-outbound-sales-2026","cat":"Blog"},{"title":"Best 5 Cekura AI Alternatives in 2026","desc":"Five Cekura AI alternatives on voice-AI passthrough, eval coverage, gateway, self-host. What each actually fixes outgrowing a voice-only testing tool.","href":"/blog/best-cekura-ai-alternatives-2026","cat":"Blog"},{"title":"MLflow LLM Tracing Alternatives in 2026: 6 LLM-Native Platforms","desc":"FutureAGI, Langfuse, Phoenix, LangSmith, Helicone, and W&B Weave as MLflow tracing alternatives in 2026 for LLM-native span trees, OTel, and evals.","href":"/blog/mlflow-llm-tracing-alternatives-2026","cat":"Blog"},{"title":"Voice AI Barge-In and Turn-Taking: A 2026 Implementation Guide","desc":"Implement barge-in and turn-taking that feels human. VAD tuning, false-barge-in defense, context preservation, and per-stage latency telemetry for 2026.","href":"/blog/voice-ai-barge-in-turn-taking-2026","cat":"Blog"},{"title":"Voice AI Observability for Vapi: A 2026 Implementation Guide","desc":"Implement voice AI observability for Vapi in 2026: native FAGI dashboard via Assistant ID, traceAI SDK path, audio_transcription and conversation rubrics.","href":"/blog/voice-ai-observability-vapi-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Scaling Claude Code in the Enterprise in 2026","desc":"Five AI gateways scored on scaling Claude Code from 50 to 5,000+ engineers: HA active-active, RBAC, 1M+ req/day audit, IdP federation.","href":"/blog/best-ai-gateways-scaling-claude-code-enterprise-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Semantic Caching in 2026","desc":"Five AI gateways for semantic caching in 2026 scored on backend choice, embedding model, TTL, per-tenant isolation, threshold tuning, hit-rate.","href":"/blog/best-ai-gateways-semantic-caching-2026","cat":"Blog"},{"title":"Best 5 Vectara Alternatives in 2026","desc":"Five Vectara alternatives on vector-store depth, hosted vs self-host, index-and-query cost. What each actually fixes outgrowing managed RAG-as-a-service.","href":"/blog/best-vectara-alternatives-2026","cat":"Blog"},{"title":"Best 5 Giskard Alternatives in 2026","desc":"Five Giskard alternatives scored on production-runtime parity, LLM-native eval depth, language reach, when a Python-only SDK is not enough.","href":"/blog/best-giskard-alternatives-2026","cat":"Blog"},{"title":"AI Agent Failure Modes in 2026: The 5-Category Taxonomy","desc":"AI agents fail in five categories, not fifty. The taxonomy production teams use to name incidents, route fixes, and close the loop.","href":"/blog/ai-agent-failure-modes-2026","cat":"Blog"},{"title":"Best 5 Laminar Alternatives in 2026","desc":"Five Laminar alternatives on OTel exporter portability, eval depth, gateway, optimizer, hosted pricing. What each actually fixes beyond agent-trace OSS.","href":"/blog/best-laminar-alternatives-2026","cat":"Blog"},{"title":"Best 5 Modal Alternatives for LLM Serving in 2026","desc":"Five Modal alternatives for production LLM serving: what each fixes vs serverless-compute design, how to keep Modal's GPU plus a real gateway and evals.","href":"/blog/best-modal-llm-serving-alternatives-2026","cat":"Blog"},{"title":"Best Multi-Agent Frameworks 2026: 7 Platforms Ranked for Production","desc":"LangGraph, CrewAI, Microsoft Agent Framework, AutoGen, Mastra, OpenAI Agents SDK, and Google ADK ranked for 2026 by debug, eval, and production readiness.","href":"/blog/best-multi-agent-frameworks-2026","cat":"Blog"},{"title":"Best 5 Ollama Alternatives for Local LLM Serving in 2026","desc":"Five Ollama alternatives scored on production scale, OpenAI-compat depth, quantization control, what each fixes when Mac-friendly local serving caps.","href":"/blog/best-ollama-local-llm-alternatives-2026","cat":"Blog"},{"title":"Deterministic LLM Evaluation Metrics (2026): The Eval Floor","desc":"Schema, regex, exact match, BLEU/ROUGE, citation-validity. Where deterministic LLM eval metrics catch 30-60 percent of failures before a judge fires.","href":"/blog/deterministic-llm-evaluation-metrics-2026","cat":"Blog"},{"title":"Future AGI vs Langfuse in 2026: Self-Improving Runtime vs Framework-Agnostic Observability","desc":"Future AGI vs Langfuse on tracing, evaluation, prompt management, deployment, security, DX. Honest verdict, May 2026 pricing, why only one closes the loop.","href":"/blog/future-agi-vs-langfuse-2026","cat":"Blog"},{"title":"Best 5 MCP Gateways for Claude Code in 2026","desc":"Five MCP gateways for Claude Code in 2026, scored on per-tool latency, server auth, tool-description scanning, session correlation, post-STDIO-RCE.","href":"/blog/best-mcp-gateways-claude-code-2026","cat":"Blog"},{"title":"Accent and Dialect Testing for Voice AI: A 2026 Methodology","desc":"Accent testing is not WER on Common Voice. 2026 methodology catches proper-noun, code-switching, filler, dialect failures production benchmarks miss.","href":"/blog/accent-dialect-testing-voice-ai-2026","cat":"Blog"},{"title":"5 Best AI Virtual Receptionist Platforms in 2026 (Tested + Ranked)","desc":"Top 5 AI virtual receptionist platforms in 2026 ranked on latency, telephony, eval depth, reliability. Honest tradeoffs plus 4 honorable mentions.","href":"/blog/best-ai-virtual-receptionist-platforms-2026","cat":"Blog"},{"title":"Best 5 Hamming Alternatives in 2026","desc":"Five Hamming alternatives on multimodal eval, gateway-runtime, OSS instrumentation, deployment. What each actually fixes when voice-only QA falls short.","href":"/blog/best-hamming-alternatives-2026","cat":"Blog"},{"title":"Best RAG Debugging Tools in 2026: 7 Platforms Compared","desc":"Phoenix, Langfuse, FutureAGI, LangSmith, Braintrust, TruLens, Galileo as the 2026 RAG debugging shortlist. Retrieval and chunk inspection.","href":"/blog/best-rag-debugging-tools-2026","cat":"Blog"},{"title":"7 Best STT Providers for Voice AI Agents in 2026 (Tested + Ranked)","desc":"Ranked STT providers for voice AI 2026: WER, real-time latency, accent and jargon handling, the rubric that scores them all on your production audio.","href":"/blog/best-stt-providers-voice-agents-2026","cat":"Blog"},{"title":"7 Best Voice Agent Monitoring Platforms in 2026","desc":"Voice agent monitoring platforms ranked for 2026 by tracing depth, named voice eval rubrics, error clustering, SLOs, inline guardrails, and audio replay.","href":"/blog/best-voice-agent-monitoring-platforms-2026","cat":"Blog"},{"title":"Future AGI vs Cekura: 2026 Voice Testing and Evaluation Comparison","desc":"Future AGI vs Cekura on voice simulation, observability, eval breadth, guardrails, optimization, deployment, compliance. Honest read, May 2026 pricing.","href":"/blog/future-agi-vs-cekura-2026","cat":"Blog"},{"title":"How to Optimize Voice Agent Latency: 12 Techniques for 2026","desc":"12 production techniques to cut voice agent latency in 2026. Streaming STT, prefix caching, prefetch tool calls, semantic cache, KV reuse, edge.","href":"/blog/how-to-optimize-voice-agent-latency-2026","cat":"Blog"},{"title":"Medical and Healthcare STT in 2026: Accent, Jargon, HIPAA","desc":"Ship clinical-grade STT in 2026: medical jargon coverage, patient accent and dialect robustness, HIPAA and BAA across audio and transcripts.","href":"/blog/medical-healthcare-stt-hipaa-2026","cat":"Blog"},{"title":"Voice AI for Healthcare and Clinical Workflows in 2026","desc":"Deploy voice AI across clinical workflows in 2026: appointment scheduling, intake, medication reminders, post-discharge follow-up under HIPAA and BAA.","href":"/blog/voice-ai-healthcare-clinical-workflows-2026","cat":"Blog"},{"title":"What Is an MCP Gateway? The 2026 Definition and Architecture Guide","desc":"Canonical 2026 MCP gateway definition: MCP 2025-11-25, OAuth 2.1, OpenTelemetry GenAI, post-April STDIO RCE threat model, 5 named approaches, 12 steps.","href":"/blog/what-is-mcp-gateway-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Retail in 2026: Personalization Routing With PII Redaction","desc":"Five AI gateways for retail in 2026 scored on personalization PII redaction, Black Friday burst load, 250ms latency, CCPA, PCI DSS v4.0.1, EU DSA scope.","href":"/blog/best-ai-gateways-retail-2026","cat":"Blog"},{"title":"Enterprise LLM Gateway for Claude Code in 2026: A Buyer's Roadmap","desc":"Staged 6-12 month roadmap for selecting an enterprise LLM gateway for Claude Code: five picks scored on pilot, expansion, procurement, TCO, exit.","href":"/blog/enterprise-llm-gateway-claude-code-2026","cat":"Blog"},{"title":"Best 6 AI Gateways for Canary Model Rollouts in 2026","desc":"Six AI gateways scored on canary rollouts: percent routing granularity, score-attached canary traffic, auto-rollback on guardrail trip, gateway gaps.","href":"/blog/best-ai-gateways-canary-model-rollouts-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Gemini CLI Multi-Model Routing in 2026","desc":"Five AI gateways scored on Gemini CLI multi-model routing 2026: 1M-context, safety-filter passthrough, Anthropic and OpenAI translation.","href":"/blog/best-ai-gateways-gemini-cli-multi-model-routing-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for RAG Pipelines in 2026","desc":"Five AI gateways for RAG pipelines in 2026: per-stage observability, embedding cost attribution, multi-vector-store routing, retrieval eval.","href":"/blog/best-ai-gateways-rag-pipelines-2026","cat":"Blog"},{"title":"Best Multi-Agent Debugging Tools in 2026: 7 Compared","desc":"FutureAGI, LangSmith, Phoenix, AgentOps, Galileo, Langfuse, Maxim as the 2026 multi-agent debugging shortlist: handoff, role, replay.","href":"/blog/best-multi-agent-debugging-tools-2026","cat":"Blog"},{"title":"Best 5 Akka SDK for LLM Alternatives in 2026","desc":"Five Akka SDK for LLM alternatives on native gateway shape, observability, runtime portability. What each actually fixes outside the Akka stack.","href":"/blog/best-akka-sdk-llm-alternatives-2026","cat":"Blog"},{"title":"The Harness Tax: What Coding-Agent Dead Weight Costs Your Engineering Org in 2026","desc":"Silent overhead burning every coding-agent budget in 2026: seat tax, context tax, tier tax, tool tax, session tax. Quantified, named, with fixes.","href":"/blog/harness-tax-coding-agent-dead-weight-2026","cat":"Blog"},{"title":"What is LLM Input/Output Validation? The 2026 Explainer","desc":"LLM input/output validation explained: schema, structure, content checks. How it differs from guardrails, what tools cover it, and how to wire it in 2026.","href":"/blog/what-is-llm-input-output-validation-2026","cat":"Blog"},{"title":"Future AGI vs Maxim Bifrost in 2026: Closed-Loop Runtime vs Go Performance","desc":"Future AGI vs Maxim Bifrost on routing, observability, cost attribution, security, deployment, DX. Honest verdict, pricing, where each falls short.","href":"/blog/future-agi-vs-maxim-bifrost-2026","cat":"Blog"},{"title":"Best LLM Tracing Tools in 2026: 6 Honest Picks","desc":"Best LLM tracing tools 2026 compared: Future AGI traceAI, Phoenix, Langfuse, OpenLLMetry, Helicone, Datadog. OTel discipline + auto-instrumentation.","href":"/blog/best-llm-tracing-tools-2026","cat":"Blog"},{"title":"Choosing an AI Gateway for Claude Code in 2026: A Complete Buyer's Guide","desc":"The 2026 buyer's guide to AI gateway for Claude Code: five picks, eight criteria, 12-week procurement, 20-question RFP, 30-day pilot.","href":"/blog/choosing-ai-gateway-for-claude-code-complete-buyers-guide-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for LLM Observability and Tracing in 2026","desc":"Five AI gateways for LLM observability and tracing 2026: OpenInference and OTel, span attributes, sampling, eval hooks, high-QPS ingestion.","href":"/blog/best-ai-gateways-llm-observability-tracing-2026","cat":"Blog"},{"title":"Best 5 AgentOps Alternatives in 2026","desc":"Five AgentOps alternatives on multi-framework instrumentation, eval pipeline, gateway, optimizer. What each actually fixes past agent-only Python tracing.","href":"/blog/best-agentops-alternatives-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for HR in 2026: Hiring, Performance, and Workforce Analytics Under EEOC Scrutiny","desc":"Five AI gateways for enterprise HR in 2026 scored on EEOC Title VII, ADA, NYC AEDT 144, Illinois AIVIA, EU AI Act, Colorado SB 24-205.","href":"/blog/best-ai-gateways-hr-2026","cat":"Blog"},{"title":"Helicone Alternatives in 2026: 6 Gateway and LLM Observability Tools","desc":"FutureAGI, Portkey, LiteLLM, Langfuse, OpenRouter, and LangSmith as Helicone alternatives in 2026 after the Mintlify acquisition. Pricing, OSS, tradeoffs.","href":"/blog/helicone-alternatives-2026","cat":"Blog"},{"title":"Best AI Gateway for Cursor Composer Multi-File Edits 2026","desc":"Six AI gateways for Cursor Composer multi-file edits in 2026, scored on semantic caching, per-developer budgets, and secret scanning at the edit boundary.","href":"/blog/best-ai-gateway-cursor-composer-multi-file-edits-2026","cat":"Blog"},{"title":"Using an MCP Gateway with Claude Code in 2026: A Practical Guide","desc":"Practical guide to using an MCP gateway with Claude Code in 2026: daily workflows, five operations with code, four production patterns, gateway picks.","href":"/blog/using-mcp-gateway-with-claude-code-practical-guide-2026","cat":"Blog"},{"title":"Best 5 DSPy Alternatives in 2026","desc":"Five DSPy alternatives scored on production runtime, optimizer breadth, instrumentation, what each fixes when research-grade prompt optimization hits prod.","href":"/blog/best-dspy-alternatives-2026","cat":"Blog"},{"title":"Best 5 OpenAI Agents SDK Alternatives in 2026","desc":"Five OpenAI Agents SDK alternatives on multi-provider routing, production patterns, observability. What each actually fixes once OpenAI-only lock-in bites.","href":"/blog/best-openai-agents-sdk-alternatives-2026","cat":"Blog"},{"title":"Best 5 Cloudflare AI Gateway Alternatives in 2026","desc":"Five Cloudflare AI Gateway alternatives scored on routing, observability depth, per-tenant chargeback, ecosystem portability, eval and optimizer loop.","href":"/blog/best-cloudflare-ai-gateway-alternatives-2026","cat":"Blog"},{"title":"Future AGI vs Kong AI Gateway in 2026: Self-Improving Runtime vs API-Gateway-Native","desc":"Future AGI vs Kong AI Gateway on routing, observability, cost attribution, security, deployment, DX. Honest verdict, pricing, self-improving loop.","href":"/blog/future-agi-vs-kong-ai-gateway-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Cline Agent Workflows in 2026","desc":"Five AI gateways scored on Cline-specific workflows 2026: per-task spend caps, tool-call observability, self-host posture, model routing.","href":"/blog/best-ai-gateways-cline-agent-workflows-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Fintech in 2026: NYDFS-Ready Routing With Model Risk Controls","desc":"Five AI gateways for fintech 2026 scored on NYDFS Part 500, revised SR 11-7 (OCC 2026-13), PCI-DSS v4.0.1, EU AI Act Annex III, DORA, and SEC 17a-4.","href":"/blog/best-ai-gateways-fintech","cat":"Blog"},{"title":"Best AI Gateway for Augment Code Workflows in 2026","desc":"Five AI gateways scored on Augment Code workflows in 2026: large-context query observability, per-dev multi-repo attribution, BYO routing, SSO/RBAC.","href":"/blog/best-ai-gateway-augment-code-workflows-2026","cat":"Blog"},{"title":"Best 5 Keywords AI Alternatives in 2026","desc":"Five Keywords AI alternatives scored on observability depth, routing intelligence, pricing above 1M req/month, optimization, deployment.","href":"/blog/best-keywords-ai-alternatives-2026","cat":"Blog"},{"title":"Best 5 Vercel AI Gateway Alternatives in 2026","desc":"Five Vercel AI Gateway alternatives scored on routing, eval/optimizer loops, enterprise RBAC, self-host posture, migration cost off the Vercel platform.","href":"/blog/best-vercel-ai-gateway-alternatives-2026","cat":"Blog"},{"title":"Self-Host LLMOps in 2026: Postgres, ClickHouse, and the Architecture Tradeoffs","desc":"Self-hosting LLM observability in 2026: Postgres vs ClickHouse, OTel collector, queue, blob storage, K8s footprint, ARM. Vendor-neutral architecture guide.","href":"/blog/llm-observability-self-hosting-guide-2026","cat":"Blog"},{"title":"Best 5 BlueJay AI Alternatives in 2026","desc":"Five BlueJay AI alternatives on scope beyond agent monitoring, self-host, gateway, optimizer. What each actually fixes beyond hosted-only eval-and-trace.","href":"/blog/best-bluejay-ai-alternatives-2026","cat":"Blog"},{"title":"Best 5 Adaline Alternatives in 2026","desc":"Five Adaline alternatives on prompt-as-config portability, provider catalog, self-host, integrated eval. What each actually fixes vs Adaline hosted-only.","href":"/blog/best-adaline-alternatives-2026","cat":"Blog"},{"title":"Best 5 Haystack Alternatives in 2026","desc":"Five Haystack alternatives on LLM-native ergonomics, pipeline portability, gateway, optimizer. What each actually fixes past a search-heritage framework.","href":"/blog/best-haystack-alternatives-2026","cat":"Blog"},{"title":"Best 5 NVIDIA NeMo Guardrails Alternatives in 2026","desc":"Five NVIDIA NeMo Guardrails alternatives on inline runtime latency, gateway, optimizer, languages. What each actually fixes outgrowing Colang flows.","href":"/blog/best-nemo-guardrails-alternatives-2026","cat":"Blog"},{"title":"Confident-AI Alternatives in 2026: 5 LLM Eval Platforms Compared","desc":"FutureAGI, Langfuse, Phoenix, Braintrust, and Galileo as Confident-AI alternatives. Pricing, OSS license, eval depth, production gaps.","href":"/blog/confident-ai-alternatives-2026","cat":"Blog"},{"title":"Best AI Gateway for Bolt.new Coding Workflows in 2026","desc":"Five AI gateways scored on Bolt.new workflows: per-project cost attribution, iteration-tree observability, error telemetry, B2B2C caps.","href":"/blog/best-ai-gateway-bolt-new-coding-workflows-2026","cat":"Blog"},{"title":"5 Best AI Voice Agent Platforms for Inbound Customer Support in 2026","desc":"Five AI voice agent platforms ranked for inbound support 2026: Vapi, Retell, ElevenLabs, LiveKit, Pipecat scored on latency, eval, reliability.","href":"/blog/best-ai-voice-agent-platforms-inbound-support-2026","cat":"Blog"},{"title":"Best 5 Cohere Platform Alternatives in 2026","desc":"Five Cohere Platform alternatives ranked on multi-provider routing, model catalog depth, embedding and rerank parity, how each frees you from Cohere-only.","href":"/blog/best-cohere-platform-alternatives-2026","cat":"Blog"},{"title":"How to Measure Voice AI Latency: The Complete 2026 Guide","desc":"Measure voice AI latency end-to-end in 2026. Per-stage budgets for STT, LLM, TTS, network. OpenInference spans, P95 SLOs, runnable traceAI code.","href":"/blog/how-to-measure-voice-ai-latency-2026","cat":"Blog"},{"title":"How to Trace Voice Agents with traceAI in 2026: STT, LLM, TTS, and Tool Spans","desc":"Trace voice agents with traceAI in 2026: STT/LLM/TTS/tool spans, OTLP transport, the FAGI Observe backend, code for LiveKit and Pipecat.","href":"/blog/traceai-opentelemetry-voice-agents-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Image and Vision LLM Routing in 2026","desc":"Five AI gateways scored on vision LLM routing in 2026: image-size limits, base64 vs URL, per-image cost, image-token estimation, streaming, guardrails.","href":"/blog/best-ai-gateways-image-vision-llm-routing-2026","cat":"Blog"},{"title":"Best 5 Datasaur Alternatives in 2026","desc":"Five Datasaur alternatives on annotation-export portability, modality coverage, self-host. What each actually fixes when NLP-annotation stops covering LLM.","href":"/blog/best-datasaur-alternatives-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Customer Support in 2026: Latency Budgets, Agent Assist, and Voice AI Passthrough","desc":"Five AI gateways for contact centers 2026 scored on sub-300ms agent-assist latency, voice-AI streaming, TCPA and ECPA, post-call audit.","href":"/blog/best-ai-gateways-customer-support-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Rate Limiting LLM Calls in 2026","desc":"Five AI gateways for rate limiting LLM calls in 2026 scored on the seven-axis rubric, provider-tier awareness, fair-share, rate-limit observability.","href":"/blog/best-ai-gateways-rate-limiting-llm-calls-2026","cat":"Blog"},{"title":"A/B Testing LLM Prompts: The Statistical Playbook (2026)","desc":"A/B testing LLM prompts without power analysis is theater. The 2026 playbook: MDE, sample sizing, matched pairs, bootstrap CIs, bandits, and rollout.","href":"/blog/ab-testing-llm-prompts-best-practices-2026","cat":"Blog"},{"title":"Self-Improving AI Agent Pipeline in 2026 (Simulate, Eval, Optimize)","desc":"Build a self-improving AI agent pipeline in 2026: synthetic users, function-call accuracy, ProTeGi rewrites. 62 to 96 percent on a refund agent.","href":"/blog/self-improving-ai-agent-pipeline","cat":"Blog"},{"title":"Who Owns Claude Code at Your Company? A Platform Team's Guide for 2026","desc":"Opinionated guide to five ownership models for Claude Code in 2026: why platform-led with security override is the default, and the 8-axis control plane.","href":"/blog/who-owns-claude-code-platform-team-guide-2026","cat":"Blog"},{"title":"Best 5 CTGT Alternatives in 2026","desc":"Five CTGT alternatives on inline guardrail latency, native gateway and eval, deployment, pricing. What each actually fixes beyond AI risk management.","href":"/blog/best-ctgt-alternatives-2026","cat":"Blog"},{"title":"What is LLM Tracing? Spans, OTel GenAI, and Sampling in 2026","desc":"LLM tracing is structured spans for prompts, tools, retrievals, and sub-agents under OTel GenAI conventions. What it is and how to implement it in 2026.","href":"/blog/what-is-llm-tracing-2026","cat":"Blog"},{"title":"Best AI Gateway for Claude Code Cost Management 2026","desc":"Six AI gateways scored on Claude Code cost management in 2026: per-developer budgets, semantic cache, audit headers, Opus-to-Haiku fallback.","href":"/blog/best-ai-gateway-claude-code-cost-management-2026","cat":"Blog"},{"title":"Best 5 BentoML Alternatives for LLM Serving in 2026","desc":"Five BentoML alternatives on LLM-native throughput, Kubernetes posture, gateway. What each actually fixes for production LLM workloads in 2026.","href":"/blog/best-bentoml-llm-alternatives-2026","cat":"Blog"},{"title":"Best 5 LLM Gateways for Monitoring Claude Code Token Spend in 2026","desc":"Five LLM gateways scored on Claude Code token-spend monitoring in 2026: session attribution, dev chargeback, dashboard slicing, alerts, BI exports.","href":"/blog/best-llm-gateways-monitoring-claude-code-token-spend-2026","cat":"Blog"},{"title":"Best 5 AgentHub Alternatives in 2026","desc":"Five AgentHub alternatives on agent portability, observability and eval depth, self-host. What each actually fixes vs a marketplace-builder product.","href":"/blog/best-agenthub-alternatives-2026","cat":"Blog"},{"title":"Intent Classification LLM Pipeline: 2026 Best Practices","desc":"A vendor-neutral 2026 intent classification pipeline. Data, judge prompt, eval, and deploy. Runs end-to-end on OpenAI + traceAI without proprietary SDKs.","href":"/blog/intent-classification-llm-pipeline-2026","cat":"Blog"},{"title":"Best 7 Open Source AI Gateways in 2026","desc":"Seven open source AI gateways for production LLM apps in 2026, ranked on license clarity, acquisition risk, supply chain, and 16 capabilities.","href":"/blog/best-open-source-ai-gateways","cat":"Blog"},{"title":"Best 5 OpenLLMetry (Traceloop) Alternatives in 2026","desc":"Five OpenLLMetry/Traceloop alternatives on instrumentation breadth, native gateway, eval and optimizer loop, community. What each actually fixes.","href":"/blog/best-openllmetry-traceloop-alternatives-2026","cat":"Blog"},{"title":"Best AI Gateway for Sourcegraph Cody Enterprise in 2026","desc":"Five AI gateways scored on Sourcegraph Cody Enterprise governance: per-developer attribution, context-retrieval cost, on-prem, embedding split.","href":"/blog/best-ai-gateway-sourcegraph-cody-enterprise-2026","cat":"Blog"},{"title":"Best 5 Langtail Alternatives in 2026","desc":"Five Langtail alternatives scored on prompt-as-API portability, deployment, eval depth, self-hosting, plus a migration plan for prompt endpoints.","href":"/blog/best-langtail-alternatives-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Logistics and Supply Chain in 2026: Route Optimization, Compliance, and Carrier Integrations","desc":"Five AI gateways for logistics and supply chain in 2026 scored on customs audit, ETA observability, carrier auth, route latency, residency, fraud.","href":"/blog/best-ai-gateways-logistics-supply-chain-2026","cat":"Blog"},{"title":"Best 5 Weights and Biases Alternatives for LLM Workflows in 2026","desc":"Five Weights and Biases alternatives scored on OpenTelemetry posture, LLM-app community depth, gateway and optimizer coverage, decoupled pricing.","href":"/blog/best-weights-biases-llm-alternatives-2026","cat":"Blog"},{"title":"Braintrust vs Datadog LLM Observability in 2026: Comparison","desc":"Braintrust vs Datadog LLM Observability in 2026: eval depth, OTel ingestion, pricing, gateway, guardrails, and the closing-the-loop axis.","href":"/blog/braintrust-vs-datadog-llm-observability-2026","cat":"Blog"},{"title":"Top Enterprise AI Gateways to Use Non-Anthropic Models in Claude Code in 2026","desc":"Five enterprise AI gateways scored on running Claude Code against non-Anthropic models in 2026: whitelist, BYO inference, audit, SOC 2 / BAA, translation.","href":"/blog/top-enterprise-ai-gateways-non-anthropic-models-claude-code-2026","cat":"Blog"},{"title":"What Is LLM Observability? A 2026 Architecture Guide","desc":"Canonical 2026 LLM observability guide: OpenInference, OTel GenAI, three pillars (cost, latency, eval), five approaches, FAGI vs Phoenix vs Langfuse.","href":"/blog/what-is-llm-observability-2026","cat":"Blog"},{"title":"Best 5 AdalFlow Alternatives in 2026","desc":"Five AdalFlow alternatives on optimizer breadth, gateway, observability, language coverage. What each actually fixes outgrowing PyTorch-style prompt libs.","href":"/blog/best-adalflow-alternatives-2026","cat":"Blog"},{"title":"Best 5 DeepEval and Confident AI Alternatives in 2026","desc":"Five DeepEval/Confident AI alternatives on eval-suite portability, runtime guardrails, dashboards, languages. What each actually fixes past the framework.","href":"/blog/best-deepeval-confident-ai-alternatives-2026","cat":"Blog"},{"title":"Best 5 Humanloop Alternatives in 2026","desc":"Five Humanloop alternatives on prompt-version portability, gateway depth, inline guardrails. What each actually fixes outgrowing prompt-engineering-first.","href":"/blog/best-humanloop-alternatives-2026","cat":"Blog"},{"title":"Logging vs LLM Observability in 2026: When Logs Stop Being Enough","desc":"What logs miss for LLM agents, what observability adds, and the 2026 tooling map across stdout, ELK, Loki, Phoenix, Langfuse, and FutureAGI.","href":"/blog/logging-vs-llm-observability-2026","cat":"Blog"},{"title":"State of LLMs at the Application Layer: 2026 Production Edition","desc":"State of frontier models, inference architecture, agents, evals, and distribution at the 2026 LLM app layer, with production picks for teams.","href":"/blog/state-of-llms-app-layer-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for Telecom in 2026: Network Operations, Customer Care, and Regulatory Compliance","desc":"Five AI gateways for telecom 2026 scored on CALEA, FCC CPNI 47 CFR 64.2001, STIR/SHAKEN, FTC TSR consent, and GDPR for EU operators.","href":"/blog/best-ai-gateways-telecom-2026","cat":"Blog"},{"title":"Best 5 Fiddler AI Alternatives in 2026","desc":"Five Fiddler AI alternatives on LLM-native instrumentation, gateway, routing, inline guardrails. What each actually fixes outgrowing ML-monitoring-first.","href":"/blog/best-fiddler-ai-alternatives-2026","cat":"Blog"},{"title":"Best AI Gateway for Windsurf Cascade Mode in 2026","desc":"Five AI gateways scored for Windsurf Cascade Mode 2026: long-session trace continuity, per-task cost, autonomous-action audit, trajectory.","href":"/blog/best-ai-gateway-windsurf-cascade-mode-2026","cat":"Blog"},{"title":"Best 5 AI Gateways for PII Redaction in LLM Calls in 2026","desc":"Five AI gateways for PII redaction in LLM calls 2026: seven-axis privacy rubric, detector latency, entity coverage, per-region policy support.","href":"/blog/best-ai-gateways-pii-redaction-llm-calls-2026","cat":"Blog"},{"title":"Best 5 AI Gateways to Route Codex CLI to Any Model in 2026","desc":"Five AI gateways scored on Codex CLI multi-provider routing in 2026: OpenAI-compatible passthrough, tool-call fidelity, cost-aware routing, latency.","href":"/blog/best-ai-gateways-codex-cli-routing-2026","cat":"Blog"},{"title":"Best 5 Apigee AI Gateway Alternatives in 2026","desc":"Five Apigee AI Gateway alternatives scored on proxy-bundle portability, eval and optimizer surfaces, ops overhead, and pricing outside Apigee.","href":"/blog/best-apigee-ai-gateway-alternatives-2026","cat":"Blog"},{"title":"Best 5 Together AI Alternatives for LLM Inference in 2026","desc":"Five Together AI alternatives on hosted inference depth, throughput, fine-tuning, VPC posture. What each actually fixes when workload outgrows the pricing.","href":"/blog/best-together-ai-alternatives-2026","cat":"Blog"},{"title":"Top 5 Enterprise AI Gateways to Track Claude Code Costs in 2026","desc":"Five enterprise AI gateways scored on Claude Code chargeback in 2026: cost-center attribution, monthly invoice reconciliation, SOX audit, finance fit.","href":"/blog/top-enterprise-ai-gateways-claude-code-cost-tracking-2026","cat":"Blog"},{"title":"Enterprise LLM Gateway for Cost Tracking Across All Coding Agents in 2026","desc":"Enterprise LLM gateway scored on cross-agent cost tracking across Claude Code, Cursor, Codex CLI, Copilot, Cline, plus FinOps export.","href":"/blog/enterprise-llm-gateway-cost-tracking-coding-agents-2026","cat":"Blog"},{"title":"Vercel AI SDK Alternatives in 2026: 5 LLM SDKs Compared","desc":"LangChain JS, Mastra, LlamaIndex.TS, OpenAI SDK, and FutureAGI as Vercel AI SDK alternatives in 2026. Pricing, OSS license, and tradeoffs.","href":"/blog/vercel-ai-sdk-alternatives-2026","cat":"Blog"},{"title":"What Is Prompt Injection? A 2026 Defense Field Guide","desc":"Definitive 2026 prompt injection field guide: direct vs indirect, OWASP LLM01:2025, MITRE ATLAS AML.T0051, MCP RCE, five defense approaches.","href":"/blog/what-is-prompt-injection-defense-2026","cat":"Blog"},{"title":"What is an LLM Dataset? Schema, Versioning, Lineage in 2026","desc":"An LLM dataset is a versioned set of input-output rows used to evaluate or fine-tune models. Schema, versioning, lineage, and 2026 tooling explained.","href":"/blog/what-is-llm-dataset-2026","cat":"Blog"},{"title":"OpenRouter Alternatives in 2026: 5 LLM Gateway Platforms Compared","desc":"Portkey, LiteLLM, TrueFoundry, Helicone, and FutureAGI as OpenRouter alternatives in 2026. Pricing, OSS license, BYOK fees, and what each won't solve.","href":"/blog/openrouter-alternatives-2026","cat":"Blog"},{"title":"Voice Agent Test Scenarios: Scale Past Manual QA in 2026","desc":"Scale voice agent testing past manual QA in 2026 with Future AGI Simulate. 4 scenario generation methods, AI test agents, CI/CD pipeline integration.","href":"/blog/voice-agent-scenario-2025","cat":"Blog"},{"title":"CrewAI vs LangGraph vs AutoGen 2026: Multi-Agent Frameworks Compared","desc":"CrewAI, LangGraph, and AutoGen compared head to head in 2026: architecture, primitives, debug, eval, and AutoGen's maintenance-mode status.","href":"/blog/crewai-vs-langgraph-vs-autogen-2026","cat":"Blog"},{"title":"LLM Tracing Best Practices in 2026: Span Hygiene, Sampling, and PII","desc":"LLM tracing best practices for 2026: OTel GenAI schema, span granularity, prompt-version tagging, tail sampling, PII redaction, cost attribution.","href":"/blog/llm-tracing-best-practices-2026","cat":"Blog"},{"title":"Future AGI Voice AI Evaluation in 2026: Latency, Tone, Audio","desc":"Future AGI voice AI evaluation in 2026: P95 latency tracking, tone scoring, audio artifact detection, refusal checks, Simulate plus Observe.","href":"/blog/future-agi-audio-level-evaluation-2025","cat":"Blog"},{"title":"LLM Benchmarks vs Production Evals in 2026: Why Public Scores Mislead","desc":"Public LLM benchmarks (MMLU, HumanEval, GSM8K) are contaminated and not predictive of production. Build domain reproductions that actually work in 2026.","href":"/blog/llm-benchmarks-vs-production-evals-2026","cat":"Blog"},{"title":"Simulated Multi-Turn LLM Evaluation: 2026 Playbook","desc":"Simulate persona x scenario x adversary, score multi-turn outcomes, gate releases. Vendor-neutral playbook with code that runs without proprietary SDKs.","href":"/blog/simulated-multi-turn-llm-evaluation-2026","cat":"Blog"},{"title":"LLM-as-Judge Best Practices in 2026: Calibration, Bias, and Cost","desc":"LLM-as-judge best practices for 2026: pick the right judge, calibrate against humans, watch length and family bias, control cost. Discipline that scales.","href":"/blog/llm-as-judge-best-practices-2026","cat":"Blog"},{"title":"W&B Weave Alternatives in 2026: 6 LLM Tracing and Eval Tools","desc":"FutureAGI, Langfuse, Phoenix, LangSmith, Braintrust, and Helicone as Weights and Biases Weave alternatives in 2026. OSS, OTel, and pricing tradeoffs.","href":"/blog/wandb-weave-alternatives-2026","cat":"Blog"},{"title":"Best AI Agent Observability Tools in 2026: 7 Honest Picks","desc":"Honest 2026 comparison of AI agent observability tools: FutureAGI, LangSmith, Langfuse, Phoenix, Braintrust, Galileo, Datadog on coverage.","href":"/blog/best-ai-agent-observability-tools-2026","cat":"Blog"},{"title":"Future AGI November 2025: Voice Persona Testing, A/B for STT-LLM-TTS","desc":"Discover Future AGI's November 2025 updates including voice agent persona testing, outbound call simulation, A/B testing for STT-LLM-TTS stacks, 30-plus.","href":"/blog/future-agi-november-roundup-2025","cat":"Blog"},{"title":"Instrument an AI Agent in Minutes with TraceAI in 2026","desc":"Instrument AI agents with TraceAI in 2026: OpenTelemetry-native Apache 2.0 spans, 20+ framework instrumentors, FITracer decorators, and 5-minute setup.","href":"/blog/instrument-your-ai-agent-traceai-2025","cat":"Blog"},{"title":"Agent Metrics Frameworks in 2026: A Decision Guide","desc":"Three agent metric frameworks own 2026: trajectory-first, task-completion-first, output-quality-first. Pick by your bug surface, not vendor pitch.","href":"/blog/agent-metrics-frameworks-2026","cat":"Blog"},{"title":"What is CrewAI? Multi-Agent Framework Explained in 2026","desc":"CrewAI is a Python framework for role-based multi-agent orchestration. Crews, agents, tasks, flows, tools, and how it differs from LangGraph and AutoGen.","href":"/blog/what-is-crewai-2026","cat":"Blog"},{"title":"What is an Agent Skill? The SKILL.md Primitive Explained for 2026","desc":"An agent skill is a folder of instructions, scripts, and resources packaged as a SKILL.md unit. What it is, how skills compose, how teams use them in 2026.","href":"/blog/what-is-agent-skill-2026","cat":"Blog"},{"title":"OpenAI AgentKit + Future AGI in 2026: Reliable Production Agents","desc":"OpenAI AgentKit (Oct 2025) + Future AGI in 2026: visual builder, traceAI auto-instrumentation, fi.evals scoring, BYOK gateway. Real code, real APIs.","href":"/blog/openai-agentkit-future-agi-2025","cat":"Blog"},{"title":"Agentic UX in 2026: Building AI-Native Interfaces (Webinar)","desc":"Webinar replay on Agentic UX in 2026 and the AG-UI protocol. Build streaming, tool-aware interfaces that work across LangGraph, CrewAI, and Mastra agents.","href":"/blog/agentic-ux-webinar-2025","cat":"Blog"},{"title":"Grafana Alternatives for LLMs in 2026: 7 Platforms Compared","desc":"FutureAGI, Datadog, Langfuse, Phoenix, Helicone, New Relic, Honeycomb as Grafana alternatives for LLM observability in 2026. Pricing, OSS, where each wins.","href":"/blog/best-grafana-alternatives-2026","cat":"Blog"},{"title":"MRR vs MAP vs NDCG: Retrieval Ranking Metrics in 2026","desc":"MRR, MAP, and NDCG decoded for 2026 retrieval and RAG systems. Worked examples, when each metric beats the others, and how to wire them into evals.","href":"/blog/what-is-mrr-map-ndcg-2026","cat":"Blog"},{"title":"Vercel AI SDK Tracing Best Practices in 2026: Edge, Streaming, OTel","desc":"Vercel AI SDK tracing best practices in 2026: experimental_telemetry, OTel GenAI, edge runtime, streaming spans, prompt versioning, and Next.js patterns.","href":"/blog/vercel-ai-sdk-tracing-best-practices-2026","cat":"Blog"},{"title":"Best Prompt Engineering Tools in 2026: 7 Platforms Compared","desc":"DSPy, FutureAGI Prompt Optimizer, PromptFoo, OpenAI Playground, Helicone Prompts, Braintrust Prompts, plus tradeoffs for 2026 prompt engineering workflows.","href":"/blog/best-prompt-engineering-tools-2026","cat":"Blog"},{"title":"Voice AI Simulation in 2026: Future AGI vs Cekura, Hamming, Bluejay, Coval","desc":"Compare voice AI simulation in 2026. Future AGI Simulate, Cekura, Hamming, Bluejay, and Coval ranked across audio evaluation, scenario generation, CI/CD.","href":"/blog/voice-ai-simulation-cekura-hamming-bluejay-coval-2025","cat":"Blog"},{"title":"Vapi vs Future AGI in 2026: Build with Vapi, Evaluate with FAGI","desc":"Vapi vs Future AGI in 2026: Vapi runs the call, FAGI evaluates it. Audio-native simulation, cross-provider benchmarking, root-cause, CI.","href":"/blog/compare-voice-ai-evaluation-vapi-vs-future-agi","cat":"Blog"},{"title":"Best Speech-to-Text APIs in 2026: Deepgram, AssemblyAI, Whisper, ElevenLabs Compared","desc":"Best STT APIs in May 2026: Deepgram Nova-3 + Flux, AssemblyAI Universal-2, Whisper, ElevenLabs Scribe v2 with WER, latency, and pricing compared.","href":"/blog/speech-to-text-apis-in-2026-benchmarks-pricing-developer-s-decision-guide","cat":"Blog"},{"title":"LLM Cost Optimization (2026): Cut Spend 30% in 90 Days","desc":"Cut LLM costs 30% in 90 days. 2026 playbook on model routing, caching, BYOK gateways, cost tracking. Includes best LLM cost-tracking tools.","href":"/blog/llm-cost-optimization-2025","cat":"Blog"},{"title":"Top Prompt Management Platforms in 2026: 7 Compared","desc":"Top prompt management platforms in 2026: Future AGI, PromptLayer, Promptfoo, Langfuse, Helicone, Braintrust, OpenAI Prompts API. Compared.","href":"/blog/top-prompt-management-platforms-2025","cat":"Blog"},{"title":"Promptfoo Alternatives in 2026: 6 LLM Eval Platforms Compared","desc":"FutureAGI, DeepEval, LangSmith, Braintrust, Phoenix, Confident-AI as Promptfoo alternatives in 2026. Pricing, OSS license, CI gating, and production gaps.","href":"/blog/best-promptfoo-alternatives-2026","cat":"Blog"},{"title":"Future AGI October 2025: OSS Stack, Vapi Integration, Scenario Testing","desc":"Future AGI's October 2025 updates: open-source AI reliability stack, Vapi voice AI integration, targeted scenario testing, and Agentic RAG.","href":"/blog/future-agi-october-roundup-2025","cat":"Blog"},{"title":"How to Debug AI Agents in 2026: Traces, Spans, and Fix Recipes","desc":"Step-by-step playbook for debugging AI agents in 2026. Real tracing decorators, span waterfall view, error propagation, tool-call diffs, and Fix Recipes.","href":"/blog/debug-ai-agents-2025","cat":"Blog"},{"title":"Best LLM Input/Output Validation Tools in 2026: 7 Compared","desc":"Pydantic AI, Instructor, Outlines, Guardrails AI, NeMo Guardrails, JSON Schema, and FutureAGI as the 2026 LLM I/O validation shortlist.","href":"/blog/best-llm-input-output-validation-tools-2026","cat":"Blog"},{"title":"Open-Source Stack for Reliable AI Agents in 2026","desc":"The 2026 OSS stack for reliable AI agents: orchestration (LangChain, LlamaIndex, Pydantic AI), gateway (LiteLLM, Open WebUI), eval and observability.","href":"/blog/open-source-stack-for-building-reliable-ai-agents","cat":"Blog"},{"title":"What is LangGraph? Stateful Agent Graphs Explained in 2026","desc":"LangGraph is LangChain's graph-based orchestration library for stateful agents. Nodes, edges, state, checkpointers, and how it differs from CrewAI.","href":"/blog/what-is-langgraph-2026","cat":"Blog"},{"title":"What is OpenRouter? The Universal LLM Marketplace Explained for 2026","desc":"OpenRouter is a hosted gateway routing one OpenAI-compatible API to 400+ models across 60+ providers, auto-fallback, unified billing.","href":"/blog/what-is-openrouter-2026","cat":"Blog"},{"title":"Build Self-Optimizing AI Agents in 2026 (Free Webinar + Guide)","desc":"Replace manual prompt tuning with eval-driven auto-optimization. 6 strategies (Bayesian, GEPA, ProTeGi), real fi.opt code, and a free 2026 webinar.","href":"/blog/agent-optimize-webinar-2025","cat":"Blog"},{"title":"Future AGI Protect 2026: Multi-Modal AI Guardrails","desc":"Future AGI Protect ships multi-modal guardrails for text, image, audio. Sub-100ms text, ~107ms image median. Toxicity, bias, privacy, prompt injection.","href":"/blog/protect-trustworthy-ai-guardrails-enterprises-2025","cat":"Blog"},{"title":"Agentic AI Evaluation 2026: A Cross-Team Framework for Reliable Agents","desc":"Agentic AI evaluation in 2026: trajectory metrics, real fi.evals code, product-engineering collaboration playbook, where Future AGI fits in the stack.","href":"/blog/agentic-ai-evaluation-2025","cat":"Blog"},{"title":"Best LLM Evaluation Tools in 2026: 7 Platforms Compared","desc":"FutureAGI, DeepEval, Langfuse, Phoenix, Braintrust, LangSmith, and Galileo as the 2026 LLM evaluation shortlist. Pricing, OSS license, and production gaps.","href":"/blog/best-llm-evaluation-tools-2026","cat":"Blog"},{"title":"Best AI Agent Orchestration Platforms in 2026: 5 Compared","desc":"LangGraph, OpenAI Agents SDK, Temporal, CrewAI, n8n compared for 2026 production agents. Code-first vs config-first vs workflow-first, honest tradeoffs.","href":"/blog/best-ai-agent-orchestration-platforms-2026","cat":"Blog"},{"title":"Best Self-Hosted LLM Observability in 2026: 7 Picks Ranked","desc":"Langfuse, Phoenix, Helicone, OpenLIT, Lunary, Comet Opik, and FutureAGI ranked on deploy footprint, scale ceiling, and self-host operational cost.","href":"/blog/best-self-hosted-llm-observability-2026","cat":"Blog"},{"title":"Athina Alternatives in 2026: 6 LLM Eval and Guardrail Platforms","desc":"FutureAGI, Langfuse, Braintrust, Phoenix, Patronus, and Helicone as Athina alternatives in 2026. Pricing, OSS license, eval-as-API, and guardrails.","href":"/blog/athina-alternatives-2026","cat":"Blog"},{"title":"UpTrain Alternatives in 2026: 7 Production-Grade Picks","desc":"FutureAGI, DeepEval, Ragas, Langfuse, Phoenix, Braintrust, and Opik as the 2026 UpTrain shortlist. License, judge depth, and self-hosting tradeoffs.","href":"/blog/uptrain-alternatives-2026","cat":"Blog"},{"title":"What is LLM Annotation? Queues, Agreement, Adjudication","desc":"LLM annotation is the human-in-the-loop labeling layer for eval datasets. Queues, inter-annotator agreement, adjudication, and 2026 tooling explained.","href":"/blog/what-is-llm-annotation-2026","cat":"Blog"},{"title":"Best AI Agent Reliability Solutions in 2026: 6 Compared","desc":"Six AI agent reliability solutions compared in 2026 across five layers: runtime guardrails, CI eval gates, span-attached scoring, clustering, closed loop.","href":"/blog/best-ai-agent-reliability-solutions-2026","cat":"Blog"},{"title":"AI Agent Cost Optimization and Observability in 2026","desc":"Agent cost optimization is an observability problem: trace-attributed cost, per-resolved-outcome, routing policies, quality-bounded swaps.","href":"/blog/ai-agent-cost-optimization-observability-2026","cat":"Blog"},{"title":"Conditional Prompts at LLM Runtime in 2026: Patterns and Pitfalls","desc":"Conditional prompt selection at runtime in 2026: routing, fallbacks, embedded conditions, version pinning, eval discipline that keeps it from drifting.","href":"/blog/conditional-prompts-llm-runtime-2026","cat":"Blog"},{"title":"Agent Evaluation Frameworks in 2026: 6 Picks Compared","desc":"Six agent eval frameworks for trajectory-first teams 2026: LangSmith, Future AGI, Braintrust, DeepEval, Phoenix, OpenAI Evals, honest tradeoffs.","href":"/blog/agent-evaluation-frameworks-2026","cat":"Blog"},{"title":"LLM Inference Performance Webinar: 2026 Update","desc":"Watch the LLM inference performance webinar, updated for 2026: continuous batching, speculative decoding, and caching that cut serving cost.","href":"/blog/inference-performance-webinar-2026","cat":"Blog"},{"title":"MLflow Alternatives in 2026: 7 LLM Eval Platforms Compared","desc":"FutureAGI, DeepEval, Langfuse, Phoenix, W&B Weave, Comet Opik, and Braintrust as MLflow alternatives for production LLM evaluation work in 2026.","href":"/blog/mlflow-alternatives-2026","cat":"Blog"},{"title":"OpenInference vs OpenLLMetry vs OpenLIT 2026: OTel for LLMs","desc":"OpenInference, OpenLLMetry, and OpenLIT compared for OpenTelemetry-based LLM observability in 2026: instrumentation, languages, semconv, and tradeoffs.","href":"/blog/openinference-vs-openllmetry-vs-openlit-2026","cat":"Blog"},{"title":"Future AGI September 2025: Agent Compass, AWS Marketplace, RBAC","desc":"See what Future AGI shipped in September 2025. Covers Agent Compass for 98 percent faster multi-agent debugging, AWS Marketplace launch, enterprise RBAC.","href":"/blog/september-update-future-agi-2025","cat":"Blog"},{"title":"Best Open Source LLM Observability in 2026: 7 Stacks Ranked","desc":"Phoenix, Langfuse, OpenLLMetry, Helicone, OpenLIT, Lunary, and FutureAGI traceAI ranked on deploy complexity, scale, OTel support, and license.","href":"/blog/best-open-source-llm-observability-2026","cat":"Blog"},{"title":"Evaluating AI Agent Skills in 2026: A Skill-Tree Playbook","desc":"Skill-level eval for agents in 2026: discrete skills, per-skill rubrics, regression sets, and CI gates. Vendor-neutral code, no proprietary SDK.","href":"/blog/evaluating-ai-agent-skills-2026","cat":"Blog"},{"title":"What is Reflection Tuning? Reflexion, Self-Refine, and 2026 Patterns","desc":"Reflection tuning is when an LLM critiques its own output and rewrites under that critique. The Reflexion / Self-Refine origins, 2026 production patterns.","href":"/blog/what-is-reflection-tuning-2026","cat":"Blog"},{"title":"LLM Benchmarks 2026: GPT-5, Claude 4.7, Gemini 2.5 Pro, Grok 4 Compared","desc":"Compare GPT-5, Claude Opus 4.7, Gemini 2.5 Pro, and Grok 4 on GPQA, SWE-bench, AIME, context, $/1M tokens, and latency. May 2026 leaderboard scores.","href":"/blog/llm-benchmarking-compare-2025","cat":"Blog"},{"title":"AI Agent Reliability Metrics 2026: Six SLOs, Not One Score","desc":"AI agent reliability is six metrics, not one composite. Task completion, tool-call success, recovery, p99 latency, guardrail trips, score, wired as SLOs.","href":"/blog/ai-agent-reliability-metrics-2026","cat":"Blog"},{"title":"LLM Fine-Tuning Guide 2026: LoRA, QLoRA, DPO, GRPO, RLHF","desc":"Fine-tune LLMs in 2026 with LoRA, QLoRA, GRPO, RLHF, DPO, IPO. Compare trl, unsloth, axolotl, DeepSpeed and learn how to evaluate fine-tuned models.","href":"/blog/llm-fine-tuning-guide-2025","cat":"Blog"},{"title":"Copilot vs Cursor vs Amazon Q Developer vs Claude Code 2026","desc":"Six AI coding agents stacked side by side: Copilot, Cursor, Amazon Q Developer, Claude Code, Codex CLI, Windsurf. Pricing, models, IDE, agent depth.","href":"/blog/github-copilot-vs-cursor-vs-codewhisperer-2025","cat":"Blog"},{"title":"Build Reliable Multi-Agent AI Flows with Future AGI in 2026","desc":"Build reliable multi-agent AI flows with Future AGI in 2026: synthetic datasets, traceAI, fi.evals, fi.simulate, Agent Command Center.","href":"/blog/build-multi-agent-ai-future-agi-2025","cat":"Blog"},{"title":"Agent Observability vs Evaluation vs Benchmarking (2026)","desc":"Observability watches. Evaluation judges. Benchmarking ranks. The conceptual map of the three terms agent teams conflate, with metrics, cadence, and tools.","href":"/blog/agent-observability-vs-evaluation-vs-benchmarking-2026","cat":"Blog"},{"title":"AI Evaluation ROI 2026: Future AGI vs In-House TCO Analysis","desc":"Future AGI vs in-house AI evaluation 2026: $400K savings, 3-year TCO breakdown, payback in weeks, build vs buy decision framework with verified pricing.","href":"/blog/ai-evaluation-roi-analysis-2025","cat":"Blog"},{"title":"RAG Evaluation Metrics in 2026: Faithfulness & More","desc":"RAG eval metrics in 2026: faithfulness, context precision, recall, groundedness, answer relevance, hallucination. With FAGI fi.evals templates.","href":"/blog/rag-evaluation-metrics-2025","cat":"Blog"},{"title":"Best Tools for Token Cost Tracking in LLMs in 2026: 7 Compared","desc":"Helicone, Langfuse, Datadog LLM cost, Braintrust, Phoenix, Portkey, FutureAGI compared on per-tenant, per-feature, per-agent token attribution.","href":"/blog/best-tools-token-cost-tracking-llms-2026","cat":"Blog"},{"title":"What is an LLM Evaluator? The 5 Types Engineering Teams Use in 2026","desc":"An LLM evaluator scores model outputs: heuristic, classifier, judge, programmatic, human. The 5 types, when each fits, and how to combine them in 2026.","href":"/blog/what-is-llm-evaluator-2026","cat":"Blog"},{"title":"What is the OpenAI Agents SDK? Loops and Handoffs in 2026","desc":"OpenAI Agents SDK is OpenAI's open-source framework for agent loops, handoffs, guardrails, and sessions. Architecture, primitives, and how to trace it.","href":"/blog/what-is-openai-agents-sdk-2026","cat":"Blog"},{"title":"Test 10,000 Voice Agent Scenarios in Minutes (2026 Guide)","desc":"Run 10,000 voice agent test scenarios in minutes in 2026 with Future AGI Simulate. Manual QA replaced by simulated callers, parallel runs, and CI/CD.","href":"/blog/voice-agent-scenarios-without-manual-qa-2026","cat":"Blog"},{"title":"What is RAG Evaluation? Frameworks, Metrics, and Gates in 2026","desc":"RAG evaluation is retrieval, generation, and end-to-end scoring under one framework. What it is, how to score each layer, which tools handle it in 2026.","href":"/blog/what-is-rag-evaluation-2026","cat":"Blog"},{"title":"How to Cut Your LLMOps Bill in 2026: 8 Concrete Levers","desc":"Eight levers to cut LLMOps spend in 2026: sampling, retention, distilled judges, semantic cache, smaller defaults, prompt caching, batches, budgets.","href":"/blog/cut-llmops-bill-2026","cat":"Blog"},{"title":"Arize AI Alternatives in 2026: 5 Honest Picks","desc":"Honest 2026 comparison of the best Arize AI alternatives: Future AGI, Langfuse, LangSmith, Braintrust, Datadog. Pricing, gateway, eval depth, license.","href":"/blog/arize-alternatives-2026","cat":"Blog"},{"title":"Future AGI August 2025: SIMULATE, Salesforce, Bedrock, Agentic RAG","desc":"Future AGI August 2025 updates: SIMULATE voice testing, function-based evals, user-level observability, Salesforce, Bedrock, Agentic RAG.","href":"/blog/august-update-future-agi-2025","cat":"Blog"},{"title":"Best AI Agent Debugging Tools in 2026: 7 Honest Picks","desc":"Honest 2026 comparison of agent debugging tools: Future AGI, LangSmith, Langfuse, Phoenix, Braintrust, Helicone, OpenLLMetry. Trace depth, replay, gaps.","href":"/blog/best-ai-agent-debugging-tools-2026","cat":"Blog"},{"title":"Best OTel Instrumentation Tools for LLMs in 2026: 6 Compared","desc":"OpenInference, traceAI, OpenLLMetry, OpenLIT, OTel-contrib, vendor SDKs as the 2026 OTel-for-LLMs shortlist. License, language, gen_ai.* support.","href":"/blog/best-otel-instrumentation-tools-llm-2026","cat":"Blog"},{"title":"AI Infrastructure Guide 2026: The Production Reference Stack","desc":"2026 reference stack for AI infrastructure: GPU compute, distributed training, MLOps, gateway routing, observability + eval, security, FinOps, real tools.","href":"/blog/ai-infrastructure-guide-2025","cat":"Blog"},{"title":"Automated Prompt Improvement in 2026: 6 Optimizers Compared","desc":"Automated prompt improvement in 2026 with six named optimizers (ProTeGi, GEPA, PromptWizard, MetaPrompt, BayesianSearch, RandomSearch) wired into CI.","href":"/blog/automated-prompt-improvement-2026","cat":"Blog"},{"title":"What is Prompt Engineering? The Practitioner's Guide for 2026","desc":"What prompt engineering means in 2026 after Bayesian, GEPA, and ProTeGi optimizers. Anatomy, techniques, tools, and where hand-tuning still earns its keep.","href":"/blog/what-is-prompt-engineering","cat":"Blog"},{"title":"TruLens Alternatives in 2026: 6 LLM Eval Platforms Compared","desc":"FutureAGI, Phoenix, Langfuse, DeepEval, Comet Opik, and Ragas as TruLens alternatives in 2026. Pricing, OSS license, feedback functions, and tradeoffs.","href":"/blog/trulens-alternatives-2026","cat":"Blog"},{"title":"Best AI Drift Detection Tools in 2026: 5 Picks by Drift Type","desc":"AI drift is five different problems wearing one name. We rank the five tools that catch them: Arize, Future AGI, Evidently, WhyLabs, Fiddler.","href":"/blog/best-ai-drift-detection-tools-2026","cat":"Blog"},{"title":"What is LiteLLM? The Universal LLM API Translator in 2026","desc":"LiteLLM, the open-source SDK and proxy giving every LLM an OpenAI-compatible API. What it is, how SDK and proxy differ, how teams use it in 2026.","href":"/blog/what-is-litellm-2026","cat":"Blog"},{"title":"Best AI Prompt Management Tools in 2026: 8 Compared","desc":"Compare 8 AI prompt management tools in 2026 across versioning, eval gates, and runtime routing. Honest tradeoffs and when to pick each.","href":"/blog/best-ai-prompt-management-tools-2026","cat":"Blog"},{"title":"AI Gateways vs LLM Gateways in 2026: 8 Platforms Compared","desc":"AI gateways govern agents, tools, MCP, voice. LLM gateways route provider calls. 8 platforms ranked across both axes with pricing and OSS license.","href":"/blog/best-ai-gateways-vs-llm-gateways-2026","cat":"Blog"},{"title":"Best LLM Annotation Tools in 2026: 8 Picked Honestly","desc":"Best LLM annotation tools in 2026 across marketplaces, self-service queues, and in-product queues. 8 platforms compared on calibration, IAA, and traces.","href":"/blog/best-llm-annotation-tools-2026","cat":"Blog"},{"title":"LLM Monitoring vs LLM Observability in 2026: A Practical Split","desc":"What LLM monitoring catches, what observability adds, where they overlap, and the 2026 tooling map across Datadog, Phoenix, Langfuse, FutureAGI.","href":"/blog/llm-monitoring-vs-llm-observability-2026","cat":"Blog"},{"title":"Real-Time LLM Evaluation in 2026: Setup, Code, Latency","desc":"Set up real-time LLM evaluation in 2026 with span-attached evals, 1 to 2 second judges, and code. 7 platforms compared, FAGI traceAI walkthrough.","href":"/blog/real-time-llm-evaluation-setup-2025","cat":"Blog"},{"title":"Voice AI Integration Guide 2026: Vapi, Retell, LiveKit, Pipecat + Eval","desc":"Voice AI integration in 2026: Vapi, Retell, LiveKit Agents, Pipecat code patterns plus traceAI instrumentation and FAGI audio evals for production.","href":"/blog/smart-voice-ai-integration-2025","cat":"Blog"},{"title":"Best LLM Monitoring Tools in 2026: 7 Platforms Compared","desc":"FutureAGI, Datadog, Langfuse, Phoenix, Helicone, Braintrust, LangSmith for LLM monitoring. Latency, drift, cost, eval pass-rate trends.","href":"/blog/best-llm-monitoring-tools-2026","cat":"Blog"},{"title":"Best AI Coding Agents 2026: 6 Picks for Senior Engineers","desc":"Best AI coding agents 2026 by job-to-be-done. Cursor, Claude Code, Cline, Aider, GitHub Copilot, Replit Agent ranked by where the agent actually lives.","href":"/blog/best-ai-coding-agents-2026","cat":"Blog"},{"title":"What is LLM Experimentation? Datasets, Runs, Variants in 2026","desc":"LLM experimentation is dataset-driven runs across prompt and model variants with attached eval scores. What it is and how to implement it in 2026.","href":"/blog/what-is-llm-experimentation-2026","cat":"Blog"},{"title":"Simulate a Voice AI Agent in 2026: A Hands-On Guide","desc":"Simulate voice AI agents in 2026 with fi.simulate.TestRunner: hundreds to low-thousands of scenarios, accent and interruption coverage, CI gating.","href":"/blog/simulate-voice-ai-agent-2025","cat":"Blog"},{"title":"Synthetic Test Data for LLM Evaluation in 2026: A Practical Guide","desc":"How to generate synthetic test data for LLM evals: contexts, evolutions, personas, contamination checks, and the OSS tools that do it well in 2026.","href":"/blog/synthetic-test-data-llm-evaluation-2026","cat":"Blog"},{"title":"Best AI Agent Governance Tools in 2026: 7 Compared","desc":"Future AGI, Credo AI, Holistic AI, Datadog, Purview, Lakera Guard, Fairly AI compared on policy authoring, audit trails, runtime enforcement for agents.","href":"/blog/best-ai-agent-governance-tools-2026","cat":"Blog"},{"title":"Future AGI + OpenAI Agents SDK: Trace + Eval in 3 Lines (2026)","desc":"Add tracing, MCP visibility, evaluations, and alerts to OpenAI Agents SDK in 3 lines with Future AGI traceAI in 2026. Apache 2.0, OpenTelemetry-native.","href":"/blog/future-agi-openai-agent-sdk-2025","cat":"Blog"},{"title":"Future AGI July 2025: OSS Eval Library, Vercel + Langfuse Tracing","desc":"Future AGI July 2025 updates: open-source eval library launch, user feedback integration, Vercel AI SDK tracing, Langfuse evaluation.","href":"/blog/july-update-future-agi-2025","cat":"Blog"},{"title":"Prompt Optimization at Scale 2026: Why Manual Tuning Fails","desc":"Manual prompt tuning fails past 50 variants. Compare Future AGI, Promptfoo, LangSmith, and Datadog for 2026 automated prompt optimization at scale.","href":"/blog/prompt-optimization-at-scale-2025","cat":"Blog"},{"title":"Best Rerankers for RAG in 2026: 7 Models Compared","desc":"Cohere Rerank 4, BGE Reranker v2-m3, Jina v2, ColBERT, Voyage rerank-2.5, mxbai, Qwen3 reranker compared on RAG-eval lift, latency, license, multilingual.","href":"/blog/best-rerankers-for-rag-2026","cat":"Blog"},{"title":"Context Engineering 2026: RAG, Memory, MCP & Evaluation","desc":"Context engineering is the production discipline around prompts in 2026. RAG, memory, MCP, tool use, evaluation, plus how Future AGI scores it.","href":"/blog/context-engineering-genai-2025","cat":"Blog"},{"title":"Future AGI vs Comet/Opik (2026): The Real Comparison","desc":"Future AGI vs Comet (Opik) in 2026. Pricing, multi-modal eval, LLM observability, G2 ratings, MLOps. Side-by-side for AI teams shipping LLM features.","href":"/blog/future-agi-vs-comet","cat":"Blog"},{"title":"Future AGI vs LangSmith 2026: LLM Eval and Observability Compared","desc":"Future AGI vs LangSmith in 2026: framework-agnostic LLM eval vs LangChain-native observability. Feature table, pricing, multimodal, verdict.","href":"/blog/future-agi-vs-langsmith","cat":"Blog"},{"title":"Future AGI vs Maxim AI in 2026: Honest Eval Comparison","desc":"Future AGI vs Maxim AI in 2026: side-by-side on eval breadth, multimodal coverage, simulation, observability, pricing, and which to pick when.","href":"/blog/future-agi-vs-maxim-ai","cat":"Blog"},{"title":"Best LLM Summarization Eval Tools in 2026: 7 Compared","desc":"DeepEval, Ragas, FutureAGI, HuggingFace Evaluate, Galileo, OpenAI Evals, Confident-AI as the 2026 summarization eval shortlist. ROUGE, BERTScore, faith.","href":"/blog/best-llm-summarization-eval-tools-2026","cat":"Blog"},{"title":"LLM Cost Tracking Best Practices in 2026: Per-User, Per-Prompt, Per-Route","desc":"LLM cost tracking 2026: token-level attribution, per-user spend caps, reasoning vs cache, gateway aggregation, drift. Practices that actually scale.","href":"/blog/llm-cost-tracking-best-practices-2026","cat":"Blog"},{"title":"Build a Generative AI Chatbot in 2026: Step-by-Step Guide","desc":"Build a generative AI chatbot in 2026: model selection, RAG, prompt-opt, evaluation, observability, guardrails, gateway. Step-by-step with current tooling.","href":"/blog/ai-chatbot-guide-2025","cat":"Blog"},{"title":"Best LLM Instrumentation Libraries in 2026: 5 Compared","desc":"OpenInference, traceAI, OpenLLMetry, OpenLIT, and Traceloop SDK as the 2026 LLM instrumentation shortlist with pip installs, code samples, and tradeoffs.","href":"/blog/best-llm-instrumentation-libraries-2026","cat":"Blog"},{"title":"Future AGI vs Braintrust in 2026: LLM Eval Platforms Compared","desc":"Future AGI vs Braintrust in 2026. Eval depth, observability, simulation, gateway, pricing, OSS status. What each platform actually does (and won't do).","href":"/blog/future-agi-vs-braintrust","cat":"Blog"},{"title":"Future AGI vs Fiddler AI 2026: Honest LLM Observability Comparison","desc":"Honest 2026 comparison of Future AGI vs Fiddler AI: LLM eval, agent observability, traditional ML monitoring, pricing, integrations, and team-fit guidance.","href":"/blog/future-agi-vs-fiddler-ai-2025","cat":"Blog"},{"title":"Future AGI vs Weights & Biases 2026: GenAI Eval vs ML Tracking","desc":"Future AGI vs Weights and Biases in 2026: GenAI evals and tracing vs experiment tracking. Verdict, head-to-head feature table, pricing, and use cases.","href":"/blog/future-agi-vs-weights-biases","cat":"Blog"},{"title":"Best LLM Evaluation Frameworks in 2026: Ranked for Production","desc":"Future AGI, DeepEval, RAGAS, Arize Phoenix, OpenAI Evals, LangSmith ranked for LLM eval in 2026. Metrics taxonomy, templates, best practices.","href":"/blog/llm-evaluation-frameworks-metrics-best-practices","cat":"Blog"},{"title":"What is an AI Gateway? Governance, Routing, and Observability in 2026","desc":"An AI gateway sits between apps and LLM providers for governance, routing, observability. What it is, how it differs from API gateways.","href":"/blog/what-is-ai-gateway-2026","cat":"Blog"},{"title":"LLM Stress Testing in 2026: Load, Adversarial, and CI Guide","desc":"How to stress-test LLMs in 2026: load testing with fi.simulate TestRunner, adversarial probes, p95 latency budgets, CI gating so failures never ship.","href":"/blog/stress-test-llm-2025","cat":"Blog"},{"title":"Top 6 AI Guardrailing Tools in 2026: Coverage, Latency, Fit","desc":"Compare top AI guardrail tools in 2026: Future AGI, NeMo Guardrails, GuardrailsAI, Lakera Guard, Protect AI, Presidio. Coverage, latency, how to choose.","href":"/blog/top-5-ai-guardrailing-tools-2025","cat":"Blog"},{"title":"Cybersecurity with GenAI Webinar (2026 Replay): Predict and Prevent","desc":"Webinar replay on cybersecurity with GenAI and intelligent agents in 2026: predictive threat detection, autonomous response, runtime guardrails for agents.","href":"/blog/cybersecurity-genai-webinar-6-2025","cat":"Blog"},{"title":"Agentic RAG in 2026: Patterns, Code, Observability","desc":"Agentic RAG in 2026: tool-using agents over vector DBs, query rewriting, multi-hop retrieval, and how to trace and evaluate every retrieve span with FAGI.","href":"/blog/agentic-rag-systems-2025","cat":"Blog"},{"title":"Future AGI vs Deepchecks 2026: LLM Eval, Pricing, G2","desc":"Future AGI vs Deepchecks in 2026: LLM evaluation, observability, prompt optimization, tabular and CV validation, pricing, G2 ratings, fit.","href":"/blog/deepchecks-vs-future-agi-2025","cat":"Blog"},{"title":"Top 5 AI Hallucination Detection Tools in 2026, Compared","desc":"The 5 best AI hallucination detection tools in 2026, ranked. Compare Future AGI, Galileo Luna, DeepEval, Phoenix, Patronus Lynx on accuracy and price.","href":"/blog/top-5-ai-hallucination-detection-tools-2025","cat":"Blog"},{"title":"How to Choose an LLM Evaluation Platform in 2026: 10 Questions","desc":"10 questions to vet any LLM evaluation platform in 2026: eval modalities, guardrails, tracing, drift, latency, scaling, and total cost of ownership.","href":"/blog/evaluation-platform-questions-2025","cat":"Blog"},{"title":"Open-Source AI Agent Stack in 2026","desc":"Open-source AI agent stack 2026: LangGraph, CrewAI, AutoGen, OpenAI Agents SDK, MS Agent Framework, Mastra, plus FAGI traceAI + ai-evaluation OSS.","href":"/blog/open-source-stack-ai-agents-2025","cat":"Blog"},{"title":"Why Enterprise AI Projects Fail in 2026: 6 Root Causes","desc":"Why so many enterprise AI projects fail in 2026: 6 root causes (KPIs, data silos, monitoring gaps, talent, technical debt, missing guardrails) and fixes.","href":"/blog/reason-enterprise-ai-project-fail-2025","cat":"Blog"},{"title":"Vibe Coding in 2026: Speed Gains, Real Risks, Production Rules","desc":"Vibe coding in 2026: prompt-driven development with Cursor, Claude Code, v0. Real productivity gains, hidden bugs, code review patterns, eval companions.","href":"/blog/vibe-coding-development-2025","cat":"Blog"},{"title":"Comparing Open-Source AI Agent Frameworks in 2026","desc":"Compare seven OSS agent frameworks for production teams in 2026, with architecture, license, maturity, latest versions, and practical tradeoffs.","href":"/blog/oss-agent-frameworks-2026","cat":"Blog"},{"title":"Best LLM Gateways in 2026: 7 Provider Routing Platforms Compared","desc":"FutureAGI ACC, Helicone, OpenRouter, Portkey, LiteLLM, Cloudflare AI Gateway, Vercel AI Gateway as 2026 LLM gateways. Routing, caching, guardrails.","href":"/blog/best-llm-gateways-2026","cat":"Blog"},{"title":"Top 10 Prompt Optimization Tools in 2026","desc":"Top 10 prompt optimization tools in 2026 ranked: FutureAGI, DSPy, TextGrad, PromptHub, PromptLayer, LangSmith, Helicone, Opik, DeepEval, Prompt Flow.","href":"/blog/top-10-prompt-optimization-tools-2025","cat":"Blog"},{"title":"Top 5 Synthetic Dataset Generators in 2026: Ranked for Production","desc":"Future AGI, Gretel, MOSTLY AI, SDV, and Snorkel ranked for synthetic dataset generation in 2026. Compare data types, privacy, agent simulation, pricing.","href":"/blog/top-5-synthetic-dataset-generators-2025","cat":"Blog"},{"title":"Voice AI Compliance in 2026: HIPAA, PCI-DSS, GDPR, EU AI Act","desc":"Voice AI regulatory compliance in 2026: HIPAA, PCI-DSS, GDPR, EU AI Act, FCC TCPA. Pre-launch audit checklist, automated testing, FAGI guardrails.","href":"/blog/voice-ai-regulatory-compliance-2026","cat":"Blog"},{"title":"What is Toolchaining? Multi-Step Tool Composition by Agents in 2026","desc":"Toolchaining is the discipline of composing multi-step tool calls in an agent: state passing, error propagation, parallel vs sequential, fine-tune line.","href":"/blog/what-is-toolchaining-2026","cat":"Blog"},{"title":"Python Decorator Tracing for LLM Apps in 2026: Patterns and Pitfalls","desc":"Decorator tracing for Python LLM apps in 2026: when to use @-tracing, when middleware fits better, OTel GenAI attributes, async pitfalls, cardinality.","href":"/blog/python-decorator-tracing-llm-2026","cat":"Blog"},{"title":"Best LLM Dataset Management Tools in 2026: 6 Compared","desc":"Best LLM dataset management tools in 2026: eval-coupled (Future AGI, Braintrust, LangSmith), annotation-first (Argilla), and generic ML (W&B, HF) compared.","href":"/blog/best-llm-dataset-management-tools-2026","cat":"Blog"},{"title":"Top 11 LLM API Providers 2026: Pricing, Latency, Context Compared","desc":"11 LLM APIs ranked for 2026: OpenAI, Anthropic, Google, Mistral, Together AI, Fireworks, Groq. Token pricing, context windows, latency, and how to choose.","href":"/blog/top-11-llm-api-providers-2025","cat":"Blog"},{"title":"AI Research Assistant Monitoring: A 2026 Playbook","desc":"Generic monitoring misses how research assistants fail. Four metrics that actually catch citation invention, source collapse, plan drift in production.","href":"/blog/ai-research-assistant-monitoring-2026","cat":"Blog"},{"title":"API vs MCP in 2026: REST/gRPC vs Model Context Protocol","desc":"API vs MCP in 2026: REST, gRPC, and GraphQL versus Model Context Protocol. Discovery, context streaming, security, versioning, and when to combine both.","href":"/blog/api-vs-mcp-difference-2025","cat":"Blog"},{"title":"Indirect Prompt Injection in 2026: XPIA, Tool Poisoning, Defense","desc":"Indirect prompt injection in 2026. Covers XPIA, tool poisoning, document-embedded prompts. FAGI Protect blocks them inline. Real defense patterns.","href":"/blog/indirect-verbal-prompts-2025","cat":"Blog"},{"title":"MarTech 2.0 GenAI Webinar (2026 Replay): Build Adaptive AI Stacks","desc":"Webinar replay on MarTech 2.0 in 2026: predictive data layers, hyper-personalization, synthetic data, adaptive agents, and the eval stack.","href":"/blog/martech-genai-webinar-5-2025","cat":"Blog"},{"title":"Prompt Injection Examples in LLMs 2026: Attacks & Defense","desc":"Real prompt injection examples in LLMs for 2026: direct, indirect, ASCII-smuggling, tool-call hijack. Ranked defense stack and working FAGI Protect code.","href":"/blog/prompt-injection-examples-llm-2025","cat":"Blog"},{"title":"Future AGI June 2025: Inline Evals, Audio Error Localizer, OSS Library","desc":"Future AGI's June 2025 updates: Inline Evaluations, Audio Error Localizer, open-source AI eval library, TypeScript ADK, Google ADK, Portkey integration.","href":"/blog/june-update-future-agi-2025","cat":"Blog"},{"title":"Future AGI + Portkey 2026: Unified Eval and Gateway","desc":"Future AGI x Portkey in 2026. Combine Portkey routing and 250+ model fallback with Future AGI traceAI eval scores. Setup in 5 minutes with Python.","href":"/blog/futureagi-portkey-integration-2025","cat":"Blog"},{"title":"Gemini 2.5 Pro in 2026: 1M Context, MCP, Deep Think, Mariner","desc":"Gemini 2.5 Pro features in May 2026: 1M token context, MCP tools, Deep Think mode, Project Mariner, Live API audio, plus how to evaluate Gemini.","href":"/blog/google-gemini-2-5-pro-2025","cat":"Blog"},{"title":"Document Summarization with LLMs in 2026: A Production Guide","desc":"Document summarization with LLMs in 2026. Extractive vs abstractive, RAG for enterprise docs, model picks, eval metrics, and a production stack.","href":"/blog/revolutionizing-document-management-llm-2025","cat":"Blog"},{"title":"Top 5 LLM Observability Tools in 2026: Ranked for Production","desc":"Future AGI, Langfuse, Arize Phoenix, Helicone, and Datadog ranked for LLM observability in 2026. Compare OTel support, eval depth, pricing, and self-host.","href":"/blog/top-5-llm-observability-tools-2025","cat":"Blog"},{"title":"What is DSPy? Stanford's Compiled Prompt Framework in 2026","desc":"DSPy is a Stanford framework that compiles LLM programs into optimized prompts. Signatures, modules, optimizers, MIPRO, and how it differs from LangChain.","href":"/blog/what-is-dspy-2026","cat":"Blog"},{"title":"Portkey Alternatives in 2026: 6 LLM Gateway and Observability Tools","desc":"FutureAGI, LiteLLM, Helicone, OpenRouter, Cloudflare AI Gateway, and Kong AI as Portkey alternatives in 2026. Pricing, OSS license, routing, tradeoffs.","href":"/blog/portkey-alternatives-2026","cat":"Blog"},{"title":"Span vs Trace in LLM Observability: What's the Difference?","desc":"A trace is one user request; a span is one operation inside that trace. OTel terminology, parent-child trees, and what makes a good LLM trace in 2026.","href":"/blog/what-is-llm-span-vs-trace-2026","cat":"Blog"},{"title":"Best Open-Source and OSS-Client LLM Eval Frameworks in 2026: 8 Test Harnesses","desc":"FutureAGI, DeepEval, Promptfoo, Ragas, UpTrain, Inspect AI, DeepChecks, MLflow Evaluate as OSS LLM eval frameworks in 2026. Compared.","href":"/blog/best-open-source-eval-frameworks-2026","cat":"Blog"},{"title":"LLM Evaluation Architecture in 2026: The Three-Tier Stack That Scales","desc":"LLM eval architecture in 2026: heuristics on every span, distilled judges on a sample, humans on the gold-set. Three-tier stack that scales.","href":"/blog/llm-evaluation-architecture-2026","cat":"Blog"},{"title":"What is AUC-ROC for LLM Evals? Operating Points and Calibration in 2026","desc":"AUC-ROC measures ranking quality of a binary classifier. Applied to LLM-judge calibration, hallucination detection, guardrail screening. When it misleads.","href":"/blog/what-is-auc-roc-llm-evals-2026","cat":"Blog"},{"title":"What is Google ADK? The Agent Development Kit Explained for 2026","desc":"Google ADK is an open-source Python, TypeScript, Go, Java framework for building and deploying agents on Vertex AI Agent Engine. What it is.","href":"/blog/what-is-google-adk-2026","cat":"Blog"},{"title":"Evaluating GenAI in Production 2026: The Full Framework","desc":"Evaluate GenAI in production in 2026: pre-deploy CI evals, online metrics, LLM-as-judge calibration, drift, safety, and how to stand up a working stack.","href":"/blog/evaluating-genai-production-2025","cat":"Blog"},{"title":"GenAI Compliance Framework 2026: EU AI Act, GDPR, CCPA","desc":"Operational GenAI compliance framework for 2026: EU AI Act phase-in, GDPR Articles 22 and 25, CCPA, HIPAA, FCRA, with evaluator-driven evidence.","href":"/blog/genai-compliance-framework-2025","cat":"Blog"},{"title":"LLM Agent Architectures in 2026: Core Components and Patterns","desc":"LLM agent architectures in 2026: ReAct, Reflexion, Plan-and-Execute, Tree-of-Thoughts, multi-agent. Memory, tools, observability with Future AGI traceAI.","href":"/blog/llm-agent-architectures-core-components","cat":"Blog"},{"title":"LLM Evaluation in 2026: Metrics, Methods, Tools, and CI","desc":"LLM evaluation in 2026: deterministic metrics, LLM-as-judge, RAG metrics, agent metrics, and how to wire offline regression plus runtime guardrails.","href":"/blog/llm-evaluation-2025","cat":"Blog"},{"title":"LLM Agents 2026: 5 Types, Applications, and Evaluation Stack","desc":"Compare 5 types of LLM agents in 2026 with real architectures, 2026 model picks (Claude 4.7, GPT-5, Gemini 3), and how to evaluate them in production.","href":"/blog/llm-agents-applications-guide-2025","cat":"Blog"},{"title":"LLM Guardrails in 2026: Implementation Guide for Safer AI","desc":"Implement LLM guardrails in 2026: 7 metrics (toxicity, PII, prompt injection), code patterns, latency budgets, and the top 5 platforms ranked.","href":"/blog/llm-guardrails-safeguarding-ai-2025","cat":"Blog"},{"title":"LLM Prompt Injection in 2026: How It Works and How to Prevent It","desc":"LLM prompt injection in 2026: direct and indirect attacks, 6 defenses (input filtering, dual LLM, output validation), and guardrail platforms.","href":"/blog/llm-prompt-injection-2025","cat":"Blog"},{"title":"Open Source vs Closed Source LLM Evaluation 2026","desc":"Pick open or closed source LLM evaluation in 2026 on cost, transparency, compliance, vendor risk, and the hybrid pattern most teams adopt.","href":"/blog/open-source-vs-closed-source-evaluations-2025","cat":"Blog"},{"title":"What is Error Analysis for LLMs? Cluster, Label, Prioritize in 2026","desc":"LLM error analysis clusters production failures, labels root causes, and prioritizes fixes. The workflow, the embeddings, and the tools teams use in 2026.","href":"/blog/what-is-error-analysis-llm-2026","cat":"Blog"},{"title":"What is LangChain? A 2026 Production Engineer's Guide","desc":"LangChain explained for 2026: what changed in v1, how LangGraph fits in, the real anatomy of the framework, production tradeoffs, and common mistakes.","href":"/blog/what-is-langchain","cat":"Blog"},{"title":"What is LLM Judge Prompting? Rubrics, Calibration, and Bias in 2026","desc":"LLM judge prompting in 2026: rubric structure, chain-of-thought, position bias, length bias, calibration, production patterns that survive real data.","href":"/blog/what-is-llm-judge-prompting-2026","cat":"Blog"},{"title":"Build a Robust MCP in 2026: Evaluate and Observe in Real Time","desc":"Build a robust MCP framework for GenAI in 2026: real-time eval, guardrails, observability, and how to wire fi.evals + traceAI to MCP servers and clients.","href":"/blog/build-mcp-evaluate-observe-2025","cat":"Blog"},{"title":"Braintrust Alternatives in 2026: 5 Honest Picks for Production AI","desc":"Honest 2026 comparison of Braintrust alternatives: Future AGI, Langfuse, Phoenix, LangSmith, Helicone. Agent trajectory eval, runtime guardrails, gateway.","href":"/blog/braintrust-alternatives-2026","cat":"Blog"},{"title":"Agentic vs Non-Agentic AI: The 2026 Definition","desc":"Agentic vs non-agentic AI explained. Workflows vs agents, the eval shape that changes, and when each architecture is the right call in production.","href":"/blog/agentic-vs-non-agentic-ai-2026","cat":"Blog"},{"title":"MCP vs A2A in 2026: Which Agent Protocol Should You Adopt?","desc":"MCP vs A2A in 2026. MCP is the Anthropic, OpenAI, Google, Microsoft backed standard. A2A is Google's peer-to-peer standard. Which to adopt and when.","href":"/blog/mcp-vs-a2a-2025","cat":"Blog"},{"title":"LLM Guardrails With Future AGI Protect in 2026: A Complete Guide","desc":"Implement LLM guardrails with Future AGI Protect in 2026. Toxicity, bias, prompt injection, data privacy. Low latency inline blocking with code samples.","href":"/blog/llm-guardrails-genai-future-agi-2025","cat":"Blog"},{"title":"Future AGI May Roundup","desc":"Future AGI May 2025 updates: MCP Server launch, 30 percent faster synthetic data generation, improved trace view with inline annotations.","href":"/blog/may-update-future-agi-2025","cat":"Blog"},{"title":"DeepEval Alternatives in 2026: 5 LLM Eval Platforms Compared","desc":"FutureAGI, Langfuse, Arize Phoenix, Braintrust, and LangSmith as DeepEval alternatives in 2026. Pricing, OSS license, eval depth, and production gaps.","href":"/blog/deepeval-alternatives-2026","cat":"Blog"},{"title":"AI Ethics Frameworks in 2026: EU AI Act + 6 Best Practices","desc":"AI ethics in 2026: six core principles, EU AI Act enforcement, OECD and NIST guidance, bias and fairness evaluation, shipping trustworthy AI in production.","href":"/blog/ethics-of-ai-framework-2025","cat":"Blog"},{"title":"What is LLM Product Analytics? A 2026 Guide","desc":"LLM product analytics: how teams join trace data to product funnels, retention, satisfaction. Tools, anatomy, mistakes, where the category is going.","href":"/blog/what-is-llm-product-analytics-2026","cat":"Blog"},{"title":"What is the Claude Agent SDK? Anthropic's Agent Loop in 2026","desc":"Claude Agent SDK is Anthropic's programmable agent harness. Python repo MIT-licensed, use under Anthropic Commercial Terms. Tools, MCP, sessions, observe.","href":"/blog/what-is-claude-agent-sdk-2026","cat":"Blog"},{"title":"Best LLM-as-Judge Platforms in 2026: 6 Compared","desc":"Future AGI, DeepEval, Galileo Luna-2, Braintrust, Phoenix, Ragas: calibrated judges, classifier cascade, deterministic floor, audit. Honest tradeoffs.","href":"/blog/best-llm-as-judge-platforms-2026","cat":"Blog"},{"title":"AI Prompting for LLMs 2026: Techniques + Examples","desc":"AI prompting techniques for 2026: zero-shot, few-shot, chain-of-thought, role, system, and how to measure prompt quality on gpt-5 and claude-opus-4-7.","href":"/blog/ai-prompting-llm-2025","cat":"Blog"},{"title":"AI LLM Test Prompts and Model Evaluation in 2026","desc":"Design AI test prompts, score model outputs, and pick a winner in 2026. Real APIs, prompt-opt loop, FAGI Evaluate, 7-step CI-ready eval pipeline.","href":"/blog/ai-llm-prompts-model-evaluation-2025","cat":"Blog"},{"title":"LLM Prompt Format 2026: 9 Patterns for GPT-5, Claude, Gemini","desc":"Nine prompt-format patterns for GPT-5, Claude Opus 4.7, and Gemini 3 workflows in 2026. Templates, eval loop, and the mistakes to avoid in production.","href":"/blog/llm-prompts-best-practices-2025","cat":"Blog"},{"title":"What is LLM Evaluation? Methods, Metrics, Tools in 2026","desc":"LLM evaluation is offline + online scoring of model outputs against rubrics, deterministic metrics, judges, and humans. Methods, metrics, and 2026 tools.","href":"/blog/what-is-llm-evaluation-2026","cat":"Blog"},{"title":"Agent Command Center: AI Gateway Control Plane (2026 Webinar)","desc":"Webinar: how routing, guardrails, and budget caps at the AI gateway layer fix the prompt injection, cost, and reliability failures teams blame on the LLM.","href":"/blog/command-centre-webinar-2026","cat":"Blog"},{"title":"Best Text-to-Speech APIs in 2026: ElevenLabs, Cartesia, Deepgram, Hume Compared","desc":"Best TTS APIs in May 2026: Cartesia Sonic 4 at 40ms, ElevenLabs v3, Deepgram Aura-2, Hume Octave, plus pricing, latency, and the right pick by use case.","href":"/blog/best-text-to-speech-providers-2026","cat":"Blog"},{"title":"Build vs Buy LLM Observability 2026: TCO, OSS Option, and the Right Call","desc":"Build vs buy LLM observability in 2026: total cost of ownership, the OSS self-host path with traceAI Apache 2.0, right call by team size and compliance.","href":"/blog/build-buy-llm-observability-2025","cat":"Blog"},{"title":"Future AGI MCP Server: Evaluate LLMs from Claude or Cursor","desc":"Run Future AGI evaluations, datasets, guardrails, and synthetic data from Claude Desktop or Cursor via MCP. Setup, code, and gotchas for 2026.","href":"/blog/future-agi-mcp-server-2025","cat":"Blog"},{"title":"Future AGI vs Confident AI in 2026: Which LLM Eval Wins?","desc":"Future AGI vs Confident AI (DeepEval) in 2026: multimodal eval, observability, OSS license, prompt-opt, and which one ships your AI app to production.","href":"/blog/future-agi-vs-confident-ai-2025","cat":"Blog"},{"title":"Modern AI Engineering in 2026: Scaling LLMs (Webinar)","desc":"Future AGI webinar with Sandeep Kaipu (Broadcom) on scaling production AI: KPI alignment, infra and data pipelines, inference cost, evaluation, guardrails.","href":"/blog/webinar-03-modern-ai-engineering","cat":"Blog"},{"title":"LLM Tool Chaining in 2026: Stop Cascading Failures in Production","desc":"LLM tool chaining in 2026. Cascading failure modes, real traceAI patterns, frameworks compared. Stop silent corruption, context loss, and timeout cascades.","href":"/blog/llm-tool-chaining-cascading-failures-production","cat":"Blog"},{"title":"Best RAG Evaluation Tools in 2026: 7 Platforms Ranked","desc":"Ragas, DeepEval, FutureAGI, Phoenix, Galileo, Langfuse, TruLens compared as the 2026 RAG eval shortlist. Faithfulness, retrieval, chunk attribution.","href":"/blog/best-rag-evaluation-tools-2026","cat":"Blog"},{"title":"How to Evaluate MCP-Connected AI Agents in Production (2026)","desc":"Evaluate MCP-connected agents in 2026: tool selection, argument correctness, task completion, OTEL tracing, and the 5-pillar production scoring framework.","href":"/blog/evaluate-mcp-connected-ai-agents-production","cat":"Blog"},{"title":"What is Ollama? The Local LLM Runtime Explained for 2026","desc":"Ollama is the open-source desktop runtime that runs Llama, Qwen, Gemma, and other open-weights LLMs locally with a one-line install. How it serves in 2026.","href":"/blog/what-is-ollama-2026","cat":"Blog"},{"title":"What is AutoGen? Microsoft's Multi-Agent Framework in 2026","desc":"AutoGen is Microsoft's open-source framework for conversational multi-agent applications. Agents, GroupChat, AgentChat, AutoGen Studio, and the v0.4 split.","href":"/blog/what-is-autogen-2026","cat":"Blog"},{"title":"Best LLM Load Testing Tools in 2026: 7 Stacks Compared","desc":"k6, Locust, vLLM benchmark, GenAI-Perf, llmperf, OpenAI Evals, FutureAGI Simulate compared on token throughput, p99, cost per test run.","href":"/blog/best-llm-load-testing-tools-2026","cat":"Blog"},{"title":"GPT-4.1 Benchmarks 2026: Should You Still Use It, or Move to GPT-5?","desc":"GPT-4.1 vs GPT-5 in 2026: SWE-bench scores, 1M token context, pricing, and the migration playbook. When to stay on 4.1 and when to switch.","href":"/blog/gpt-4-1-benchmarks-2025","cat":"Blog"},{"title":"LLM Observability and Monitoring in 2026: The Field Guide","desc":"What LLM observability means in 2026: traces, spans, evals, span-attached scores. Top 5 platforms compared, real traceAI code, alerts.","href":"/blog/llm-observability-monitoring-2025","cat":"Blog"},{"title":"Future AGI April 2025: Compare Data, Audio Evals, OpenAI Agents SDK","desc":"Future AGI April 2025 updates: Compare Data for LLM comparison, Knowledge Base synthetic data, Audio Evaluations, OpenAI Agents SDK.","href":"/blog/april-update-future-agi-2025","cat":"Blog"},{"title":"Mistral Small 3.1 in 2026: Benchmarks, Lineup, Comparison","desc":"Mistral Small 3.1 in May 2026: 128k context, vision, 80.6% MMLU, Apache 2.0. Plus where Small 3.2, Medium 3, and Mistral Large 2 fit the lineup.","href":"/blog/mistral-small-3-1-2025","cat":"Blog"},{"title":"Top 5 LLM Evaluation Tools 2026: Future AGI, Galileo, Arize Compared","desc":"The 5 LLM evaluation tools worth shortlisting in 2026: Future AGI, Galileo, Arize AI, MLflow, Patronus. Features, pricing, and which workload each wins.","href":"/blog/top-5-llm-evaluation-tools-2025","cat":"Blog"},{"title":"Gemini 2.5 Pro in 2026: Is It Still Worth Using After Gemini 3.1 Pro?","desc":"Gemini 2.5 Pro in May 2026: pricing, benchmarks, retirement status, and whether to upgrade to Gemini 3.1 Pro for new builds. With migration checklist.","href":"/blog/gemini-2-5-pro-2025","cat":"Blog"},{"title":"How to Cut RAG Hallucinations in 2026: Future AGI Playbook","desc":"Cut RAG hallucinations 2026 with the Future AGI eval loop. Context Adherence and Groundedness, real fi.evals code, chunk and retriever and reranker tuning.","href":"/blog/rag-hallucinations-future-agi-2025","cat":"Blog"},{"title":"ROI of AI Explainability Tools in 2026: SHAP, LIME, Captum","desc":"Measure ROI of AI explainability tools in 2026: SHAP, LIME, Captum, Alibi, TransformerLens, KPIs, finance and healthcare results, real audit savings.","href":"/blog/roi-ai-explainability-2025","cat":"Blog"},{"title":"What is RAG Fluency? Distinct from Groundedness, Measured in 2026","desc":"RAG fluency scores how well a generated answer reads, distinct from groundedness, accuracy, relevance. What it is, measurement, when it matters.","href":"/blog/what-is-rag-fluency-2026","cat":"Blog"},{"title":"AI Compliance Guardrails for Enterprise LLMs (May 2026 Guide)","desc":"Map enterprise LLMs to GDPR, EU AI Act and NIST AI RMF in 2026: input/output guardrails, bias audits, explainability, and a real FAGI Protect setup.","href":"/blog/ai-compliance-guardrails-enterprise-llms-2025","cat":"Blog"},{"title":"Real-Time vs Batch LLM Monitoring in 2026: A Decision Framework","desc":"When real-time LLM evaluation beats batch, when batch wins, and the cost-and-latency tradeoffs across guardrails, judge sampling, and offline evals.","href":"/blog/real-time-vs-batch-llm-monitoring-2026","cat":"Blog"},{"title":"Langfuse vs LangSmith 2026: Head-to-Head LLM Observability","desc":"Langfuse vs LangSmith 2026 head-to-head: license, framework neutrality, prompts, datasets, eval, self-host, the unified-stack axis.","href":"/blog/langfuse-vs-langsmith-2026","cat":"Blog"},{"title":"Best LLM Judge Models in 2026: 8 Models Ranked","desc":"Eight LLM judge models compared on human correlation, cost per score, latency, and self-preference bias. Pick by your rubric, not by SummEval.","href":"/blog/best-llm-judge-models-2026","cat":"Blog"},{"title":"Chain of Draft Prompting in 2026: Cut Tokens 80%, Match CoT Accuracy","desc":"Chain of Draft (Xu et al. 2025) cuts reasoning tokens by ~80% while matching Chain of Thought accuracy on math, symbolic, and commonsense benchmarks.","href":"/blog/chain-of-draft-llm-2025","cat":"Blog"},{"title":"Manus AI in 2026: Pricing, GAIA Scores, and the Best Alternatives","desc":"Manus AI in May 2026: current pricing, GAIA Level 3, agent quality, and how it compares to Devin, Cursor, Replit Agent, Claude Code, and Operator.","href":"/blog/manus-ai-comparison-2025","cat":"Blog"},{"title":"Best LLM Routers and Load Balancers in 2026: 7 Compared","desc":"OpenRouter, Portkey, LiteLLM, RouteLLM, Martian, FutureAGI, Kong AI for LLM routing in 2026. Compared on routing depth, fallbacks, and pricing.","href":"/blog/best-llm-routers-load-balancers-2026","cat":"Blog"},{"title":"Future AGI vs Arize AI 2026: Best LLM Eval Tool?","desc":"Future AGI vs Arize AI in 2026: eval coverage, traceAI vs Phoenix OSS, multimodal eval, agent simulation, gateway, and pricing for production teams.","href":"/blog/future-agi-arize-ai-llm-evaluation-2025","cat":"Blog"},{"title":"Build an LLM Evaluation Framework in 2026: Code & Metrics","desc":"Build an LLM evaluation framework from scratch in 2026. Deterministic, rubric, LLM-as-judge, and agent evals, with working Python code and a CI gate.","href":"/blog/build-llm-evaluation-framework-2025","cat":"Blog"},{"title":"LLM Guardrails Deployment in 2026: Patterns + Real Code","desc":"Deploy LLM guardrails in 2026 with sub-2s inline checks, defensive layers, fallbacks, monitoring. Real Future AGI code, EU AI Act deadlines, 5 steps.","href":"/blog/llm-gaurdrails-deployement-2025","cat":"Blog"},{"title":"LLM Observability in 2026: A CTO Playbook for Tools and Tradeoffs","desc":"LLM observability in 2026 for CTOs. Metrics, logs, traces, tool selection, lifecycle integration, an Instacart case study, plus traceAI in production.","href":"/blog/llm-observability-transparency-2025","cat":"Blog"},{"title":"Best LLM Experimentation Tools in 2026: 5 Stats-Grade Picks","desc":"Five LLM experimentation tools ranked for 2026 on what actually ships a winning prompt: paired evals, bootstrap CIs, span-attached scoring, and CI gates.","href":"/blog/best-llm-experimentation-tools-2026","cat":"Blog"},{"title":"Patronus Alternatives in 2026: 6 LLM Eval and Agent Platforms","desc":"FutureAGI, Langfuse, Braintrust, Phoenix, DeepEval, Helicone as Patronus alternatives in 2026. Pricing, OSS license, hallucination detection, eval.","href":"/blog/patronus-alternatives-2026","cat":"Blog"},{"title":"What is Haystack? Deepset's RAG and Agents Framework in 2026","desc":"Haystack is Deepset's open-source pipeline framework for RAG and agents. Components, pipelines, document stores, agents, and the Haystack 2.x rewrite.","href":"/blog/what-is-haystack-2026","cat":"Blog"},{"title":"Top Agentic AI Frameworks in 2026: 7 Picks Compared","desc":"Compare the top agentic AI frameworks in 2026: LangGraph, OpenAI Agents SDK, Microsoft Agent Framework, CrewAI, AutoGen, Mastra, and PydanticAI.","href":"/blog/agentic-ai-frameworks-2025","cat":"Blog"},{"title":"Agentic AI vs Generative AI (2026): Differences & ROI","desc":"Agentic AI vs generative AI in 2026. Real differences, when to pick each, how to combine them, and how to evaluate both for production ROI.","href":"/blog/agentic-ai-vs-generative-ai-2025","cat":"Blog"},{"title":"Grok 4 vs Grok 3 Review 2026: Benchmarks, 2M Context, Tool Use","desc":"Grok 4, Grok 4.1 Fast, and Grok 4.3 reviewed for 2026. Covers AIME, GPQA, HLE scores, 256K vs 2M context, $0.20/1M pricing, and where Grok 3 fits today.","href":"/blog/grok-3-technical-review-2025","cat":"Blog"},{"title":"LLM Inference in 2026: How It Works, Latency & Cost","desc":"How LLM inference works in 2026: tokenization, KV cache, decoding, latency targets (TTFT under 500ms), cost math, and 7 optimizations that move the needle.","href":"/blog/llm-inference-human-prompts-2025","cat":"Blog"},{"title":"Multi-Agent AI Systems in 2026: Frameworks, Patterns, Production","desc":"Multi-agent AI systems in 2026: CrewAI, LangGraph, AutoGen, OpenAI Agents SDK, MS Agent Framework compared. Patterns, traceAI observability, eval, gateway.","href":"/blog/multi-agent-systems-2025","cat":"Blog"},{"title":"Vector Databases and Knowledge Graphs for RAG in 2026","desc":"Vector databases vs knowledge graphs for RAG in 2026. Pinecone, Weaviate, Qdrant, Milvus, Chroma vs Neo4j, GraphRAG, LightRAG. Decision matrix.","href":"/blog/vector-databases-knowledge-graphs-rag-2025","cat":"Blog"},{"title":"LLM Reasoning in 2026: o3, GPT-5, Claude 4.7, DeepSeek R1 Guide","desc":"How LLM reasoning works in 2026: o3, GPT-5 thinking, Claude 4.7 extended thinking, DeepSeek R1, chain-of-thought, tree-of-thoughts.","href":"/blog/llm-reasoning-2025","cat":"Blog"},{"title":"Best LLM Chatbot Evaluation Tools in 2026: 7 Compared","desc":"DeepEval, FutureAGI, Confident-AI, Galileo, Coval, Langfuse, and Maxim as the 2026 chatbot eval shortlist. Multi-turn, persona, escalation, satisfaction.","href":"/blog/best-llm-chatbot-evaluation-tools-2026","cat":"Blog"},{"title":"Evaluating AI With Confidence in 2026: A Practical Guide","desc":"Evaluate AI with confidence in 2026. Early-stage evals, multi-modal scoring, custom metrics, error localization, FAGI workflow, and CI patterns that ship.","href":"/blog/evaluating-ai-with-confidence","cat":"Blog"},{"title":"Model Context Protocol (MCP) in 2026: Standard for AI Tool Use","desc":"MCP became the de facto AI tool-use standard in 2025-2026: Anthropic, OpenAI, and Google all adopted it. Architecture, SDKs, security, gateway options.","href":"/blog/model-context-protocol-mcp-2025","cat":"Blog"},{"title":"CI/CD for AI Agents in 2026: Eval Gates, Regression Suites, Canary Rollouts","desc":"CI/CD pipelines for AI agents in 2026: eval gates, golden datasets, canary deploys, regression suites. GitHub Actions and GitLab patterns that ship safely.","href":"/blog/ci-cd-for-ai-agents-best-practices-2026","cat":"Blog"},{"title":"G-Eval vs DeepEval Metrics in 2026: Where Each Fits","desc":"G-Eval rubric-based LLM judges vs DeepEval's full metric suite, how they differ, and where FutureAGI Turing eval models fit alongside both in 2026.","href":"/blog/g-eval-vs-deepeval-metrics-2026","cat":"Blog"},{"title":"Future AGI vs Galileo AI in 2026: Honest Comparison","desc":"Future AGI vs Galileo AI for LLM evaluation in 2026: Apache 2.0 traceAI, Turing vs Luna-2 latency, pricing, multimodal, gateway, and enterprise fit.","href":"/blog/future-agi-galileo-ai-llm-evaluation-2025","cat":"Blog"},{"title":"How Multimodal LLMs Work in 2026: Vision Encoders, Fusion, and Decoders","desc":"Multimodal LLM internals in 2026. Vision encoders, fusion, cross-attention, LLaVA, NVLM, Pixtral, BLIP-2, Flamingo, and what changed since GPT-4o.","href":"/blog/exploring-how-multimodal-large-language-models-work","cat":"Blog"},{"title":"LLM Application Tech Stack in 2026: Layer-by-Layer Guide","desc":"The complete 2026 LLM application stack: foundation models, orchestration, vector DBs, LLMOps, gateways. Compare every layer with the leaders in each.","href":"/blog/llm-application-tech-stack-2025","cat":"Blog"},{"title":"Best Tools to Monitor Multi-Agent Systems in 2026: 7 Compared","desc":"Galileo Agent Graph, Maxim, AgentOps, LangGraph Studio, Arize, FutureAGI, Phoenix on handoff metrics and parallel-step analysis. Compared.","href":"/blog/best-tools-monitoring-multi-agent-systems-2026","cat":"Blog"},{"title":"What is LlamaIndex? RAG and Agents Framework in 2026","desc":"LlamaIndex is the open-source data framework for RAG and agents over enterprise data. Indexes, query engines, agents, workflows, and 0.14 architecture.","href":"/blog/what-is-llamaindex-2026","cat":"Blog"},{"title":"Best LLM Agent Memory Tools in 2026: 6 Honest Picks","desc":"Mem0, Zep, Letta, LangMem, MotorHead, and Postgres+pgvector for LLM agent memory in 2026. Honest tradeoffs on recall, freshness, and contradictions.","href":"/blog/best-llm-agent-memory-tools-2026","cat":"Blog"},{"title":"AI Guardrail Metrics 2026: Accuracy, Bias, Safety, PII","desc":"The 8 guardrail metrics every production LLM team tracks in 2026: PII, jailbreak, toxicity, bias, faithfulness, latency, refusal rate, drift. With tooling.","href":"/blog/ai-guardrail-metrics","cat":"Blog"},{"title":"Best Retrieval Quality Monitoring Tools in 2026: 7 Compared","desc":"Phoenix, Galileo, FutureAGI, Langfuse, Ragas, TruLens, UpTrain as 2026 retrieval quality monitoring picks. Recall@k, faithfulness, context relevance.","href":"/blog/best-retrieval-quality-monitoring-tools-2026","cat":"Blog"},{"title":"ChatGPT Jailbreak in 2026: How It Works & Defenses","desc":"ChatGPT jailbreak in 2026: DAN family, prompt injection, role-play, encoded payloads, and how FAGI Protect blocks them as a runtime guardrail layer.","href":"/blog/jailbreaking-chatgpt-2025","cat":"Blog"},{"title":"What Is RAG (Retrieval-Augmented Generation)? 2026 Guide for LLM Teams","desc":"Retrieval-Augmented Generation for LLMs in 2026: how it works, hybrid plus reranker stack, eval metrics, FAGI companion for production.","href":"/blog/understanding-rag-llm-a-powerful-approach-for-ai-models","cat":"Blog"},{"title":"Vellum Alternatives in 2026: 6 LLM Eval and Agent Platforms Compared","desc":"FutureAGI, Braintrust, Langfuse, LangSmith, Phoenix, and Helicone as Vellum alternatives in 2026. Pricing, OSS license, eval depth, and tradeoffs.","href":"/blog/vellum-alternatives-2026","cat":"Blog"},{"title":"What is RAG Observability? Tracing Retrieval in 2026","desc":"RAG observability is span-level tracing of retrieval, reranking, and generation, with chunk-level scores and grounding metrics. What it is, how to ship.","href":"/blog/what-is-rag-observability-2026","cat":"Blog"},{"title":"MLOps vs LLMOps in 2026: What Actually Changed","desc":"MLOps vs LLMOps in 2026. Where the practices overlap, where they diverge, and how the LLM stack reshapes training, eval, monitoring, and deployment.","href":"/blog/mlops-vs-llmops-2026","cat":"Blog"},{"title":"Detect Hallucinations in Generative AI: 6 Methods That Work in 2026","desc":"Detect AI hallucinations in production in 2026: ChainPoll, NLI, SelfCheckGPT, RAG faithfulness, FAGI eval, and human review. Code, latency, and trade-offs.","href":"/blog/detect-hallucination-generative-ai-2025","cat":"Blog"},{"title":"How to Evaluate RAG Systems in 2026: Metrics, Methods, Tools","desc":"How to evaluate RAG systems in 2026: retrieval, faithfulness, hallucination, chunk attribution, query coverage metrics, tool comparison, Future AGI fit.","href":"/blog/evaluating-rag-systems-ensuring-your-llm-remembers-what-it-reads","cat":"Blog"},{"title":"LLMOps 2026: Monitor, Optimize, and Secure Production LLMs","desc":"Monitor, optimize, and secure LLMs in production in 2026. Three pillars of observability, ethical guardrails, root cause analysis, and the tools that ship.","href":"/blog/llmops-secrets-how-to-monitor-optimize-llms-for-speed-security-accuracy","cat":"Blog"},{"title":"LLM Safety and Compliance Guide for 2026: A Practical Playbook","desc":"EU AI Act, NIST AI RMF, ISO 42001, jailbreaks, PII, and hallucination gates: a 2026 LLM safety playbook for production teams shipping under regulation.","href":"/blog/llm-safety-compliance-guide-2026","cat":"Blog"},{"title":"Multi-Turn LLM Evaluation in 2026: A Practical Guide","desc":"What multi-turn LLM evaluation actually measures in 2026, why single-turn metrics fail on agents, and the OSS and commercial tools that handle it.","href":"/blog/multi-turn-llm-evaluation-2026","cat":"Blog"},{"title":"AI Chatbot Build Guide 2026: RAG, Evals, Guardrails","desc":"End-to-end 2026 guide for building production AI chatbots: model picks, RAG, hallucination evals, traceAI observability, and runtime guardrails.","href":"/blog/ai-chatbot-guide-future-agi-2025","cat":"Blog"},{"title":"Galileo Alternatives in 2026: 7 Honest Picks for Eval Teams","desc":"Honest 2026 comparison of Galileo alternatives: Future AGI, LangSmith, Langfuse, Phoenix, Braintrust, Helicone, Datadog. Eval, gateway, Luna-2 cost.","href":"/blog/galileo-alternatives-2026","cat":"Blog"},{"title":"AI Failures and Smart Evaluation in 2026: Webinar Replay","desc":"Watch the Future AGI webinar on AI evaluation, updated for 2026. Covers why classic test suites miss agent failures and a live evals walkthrough.","href":"/blog/webinar-01-ai-failures-smart-evaluation-techniques","cat":"Blog"},{"title":"What is Prompt Versioning? Registries, Labels, and Rollback in 2026","desc":"Prompt versioning treats prompts as code: unique ids, environment labels, eval-gated rollouts, one-call rollback. What it is and how to ship it in 2026.","href":"/blog/what-is-prompt-versioning-2026","cat":"Blog"},{"title":"Synthetic Data Generation for Bias Mitigation in 2026","desc":"How synthetic data generation closes bias in AI training in 2026: five methods, fairness audits, the closed-loop workflow with FAGI Dataset + Fairness.","href":"/blog/synthetic-data-generation-bias-2025","cat":"Blog"},{"title":"LLM Testing in Production: The 2026 Practitioner's Playbook","desc":"A 6-step LLM testing loop for 2026: instrument with OTel, score spans, gate releases in CI, simulate, sample live traffic, optimize prompts on failures.","href":"/blog/llm-testing-playbook-2026","cat":"Blog"},{"title":"What is vLLM? The High-Throughput LLM Serving Engine in 2026","desc":"vLLM is the open-source LLM serving engine that pioneered PagedAttention and continuous batching. How it serves, how teams use it in production 2026.","href":"/blog/what-is-vllm-2026","cat":"Blog"},{"title":"Multimodal AI in 2026: GPT-5, Claude Opus 4.7, Gemini 2.5 Pro Picked Apart","desc":"Multimodal AI in May 2026: how GPT-5, Claude Opus 4.7, and Gemini 2.5 Pro handle text plus image plus audio plus video. Real production patterns.","href":"/blog/multimodal-ai-2025","cat":"Blog"},{"title":"LangChain Callbacks 2026: Handlers, Events, Tracing Guide","desc":"LangChain callbacks in 2026: every lifecycle event, sync vs async handlers, runnable config patterns, and how to wire callbacks into OpenTelemetry traces.","href":"/blog/understanding-langchain-callback-how-to-use-it-effectively","cat":"Blog"},{"title":"Custom LLM Eval Metrics (2026): The Three-Part Contract That Works","desc":"Custom LLM eval metrics in 2026: a tight criterion, a calibrated corpus, a stability check. Patterns, code, and pitfalls, all in one guide.","href":"/blog/custom-llm-eval-metrics-best-practices-2026","cat":"Blog"},{"title":"AI Chatbot Development in 2026: LLMs, RAG, Agentic Techniques","desc":"How to build production AI chatbots in 2026. Compare GPT-5, Claude Opus 4.7, Gemini 3, Llama 4. RAG, agentic memory, eval, and handoff patterns that ship.","href":"/blog/developing-smarter-chatbots-essential-ai-chatbot-development-techniques-for-2025","cat":"Blog"},{"title":"AI Fairness in 2026: Detect & Fix LLM Bias (Real Code)","desc":"Detect demographic parity, equal opportunity, and toxicity bias in LLM outputs in 2026. Real code with Future AGI evals + guardrails, EU AI Act deadlines.","href":"/blog/fairness-in-ai-how-to-detect-and-mitigate-bias-in-llm-outputs-using-future-agi-metrics","cat":"Blog"},{"title":"LangChain QA Evaluation in 2026: Metrics, Patterns, Tools","desc":"Evaluate LangChain QA chains in 2026: metrics, golden datasets, LangSmith vs LangChain evaluators vs Future AGI, and a working code walkthrough.","href":"/blog/langchain-qa-evaluation-2025","cat":"Blog"},{"title":"What is an MCP Server? Architecture, Transports, and Primitives in 2026","desc":"An MCP server exposes tools, resources, and prompts to LLM clients via Model Context Protocol. Architecture, transports (stdio, SSE, HTTP), lifecycle.","href":"/blog/what-is-mcp-server-2026","cat":"Blog"},{"title":"Llama vs Traditional AI Models (2026): Llama 4 vs GPT, BERT","desc":"Llama 4 vs traditional AI models in 2026. Open-source vs proprietary, architecture, efficiency, customization, and how to evaluate LLM outputs.","href":"/blog/llama-traditional-ai-models-2025","cat":"Blog"},{"title":"Multimodal Image-to-Text Models in 2026: GPT-5o, Claude 4.7, Gemini 3","desc":"Compare GPT-5o, Claude Opus 4.7, Gemini 3 Pro, and Llama 4 vision in 2026. Covers MMMU, MathVista, MMVet benchmarks plus eval and tracing patterns.","href":"/blog/the-future-of-ai-advancements-in-multimodal-image-to-text-models","cat":"Blog"},{"title":"Best No-Code LLM Builders in 2026: 7 Drag-and-Drop Platforms","desc":"Dify, Flowise, Langflow, n8n, Vapi, Voiceflow, Stack AI for no-code LLM apps in 2026. Compared on visual builders, agents, voice, OSS license, and pricing.","href":"/blog/best-no-code-llm-builders-2026","cat":"Blog"},{"title":"Prompt Injection in 2026: Attack Types and How to Defend","desc":"Prompt injection in 2026: direct, indirect, jailbreak, and covert attacks explained, plus a working defense pattern with the FAGI Protect Guardrails SDK.","href":"/blog/prompt-injection-2025","cat":"Blog"},{"title":"Vector Chunking 2026: Strategies, Sizes, and Retrieval Wins","desc":"Vector chunking in 2026: fixed, semantic, late, hierarchical, agentic, and SPLADE-style sparse chunking compared with sizes, retrieval gains, and pitfalls.","href":"/blog/vector-chunking-2025","cat":"Blog"},{"title":"Transformer Architecture Evaluation 2026: Metrics, Benchmarks, Tradeoffs","desc":"Evaluate transformer architectures in 2026: attention quality, perplexity, MMLU, GLUE, SQuAD, HellaSwag, training stability, inference throughput, checks.","href":"/blog/evaluating-transformer-architectures-key-metrics-and-performance-benchmarks","cat":"Blog"},{"title":"Controllable TalkNet in 2026: TTS Guide and Setup","desc":"Controllable TalkNet in 2026: how the TTS model works, pitch and duration controls, where to actually run it, ethics, and evaluating voice output.","href":"/blog/talknet-hugging-face-2025","cat":"Blog"},{"title":"How to Implement Voice AI Observability in 2026: Full Guide","desc":"Implement voice AI observability in 2026 for Vapi, Retell, LiveKit, Pipecat. Real traceAI code, latency SLOs, audio metrics, live eval scoring.","href":"/blog/implement-voice-ai-observability-2026","cat":"Blog"},{"title":"LLM Leaderboard Explained 2026: Arena, MMLU, GPQA, SWE-bench","desc":"How LLM leaderboards work in 2026: Chatbot Arena, MMLU, MMMU, GPQA, SWE-bench, HumanEval. Current top models and how to evaluate them on your own data.","href":"/blog/llm-leaderboard-explained","cat":"Blog"},{"title":"Future AGI Prompt Optimize 2026: 6 Algorithms, Code Inside","desc":"Future AGI Prompt Optimize in 2026: six search algorithms (BayesianSearch, MIPRO, GEPA, ProTeGi, PromptWizard, Random) with code, evals, and CI gating.","href":"/blog/prompt-optimization-future-agi-2025","cat":"Blog"},{"title":"Ragas Alternatives in 2026: 7 Production RAG Eval Picks","desc":"FutureAGI, DeepEval, TruLens, Phoenix, Langfuse, Galileo, Braintrust as the 2026 Ragas shortlist. Faithfulness, retrieval, production gaps compared.","href":"/blog/ragas-alternatives-2026","cat":"Blog"},{"title":"DeepSeek R1 vs GPT-5, Claude 4.7, Gemini 3 Pro (2026)","desc":"DeepSeek R1 and V3 compared to GPT-5, Claude Opus 4.7, Gemini 3 Pro in 2026: architecture, benchmarks, cost, and how to evaluate on workload.","href":"/blog/evaluating-deepseek-ai-vs-top-competitors","cat":"Blog"},{"title":"OpenAI Operator in 2026: GPT-5, ChatGPT Atlas, and Alternatives","desc":"OpenAI Operator in 2026: how it folded into GPT-5 and ChatGPT Atlas, what it can do, plus 6 alternatives compared (Claude, Browserbase, Hyperbrowser).","href":"/blog/openai-operator-2025","cat":"Blog"},{"title":"Validate Synthetic Datasets With Future AGI in 2026 (5 Steps)","desc":"Validate synthetic datasets with Future AGI in 2026. Five step workflow covering ingest, quality, bias, real vs synthetic, and observability with code.","href":"/blog/validate-synthetic-data-with-future-agi-2025","cat":"Blog"},{"title":"Generative AI Trends 2026: 8 Shifts Reshaping Builds & Buys","desc":"Eight 2026 generative AI trends: agentic AI, multimodal, GPT-5, Claude 4.7, Gemini 2.5 Pro, on-device, MCP, evals, gateways, with budgets.","href":"/blog/generative-ai-trends-2026","cat":"Blog"},{"title":"LangChain RAG Observability in 2026: traceAI + Eval Stack","desc":"Trace and evaluate every LangChain RAG step in 2026 with Future AGI traceAI-langchain. Compare recursive, semantic, CoT retrieval with grounded metrics.","href":"/blog/langchain-rag-observability-2025","cat":"Blog"},{"title":"Voice AI Evaluation Infrastructure 2026: A Developer Guide","desc":"Voice AI evaluation infrastructure in 2026: five testing layers, STT/LLM/TTS metrics, synthetic harness, traceAI, and FAGI Simulate.","href":"/blog/voice-ai-evaluation-infrastructure-developers-guide","cat":"Blog"},{"title":"AI Red Teaming for GenAI in 2026: Tools, Attacks, Playbook","desc":"AI red teaming for generative AI in 2026: 5 attack categories, top tools (Future AGI Protect, Garak, PyRIT, Lakera), CI playbook, and how to score risk.","href":"/blog/ai-red-teaming-genai-2025","cat":"Blog"},{"title":"Chain of Thought Prompting in 2026: Guide for GPT-5 + Claude 4.7","desc":"Chain of thought prompting in 2026: how CoT works in GPT-5, Claude 4.7 extended thinking, and DeepSeek R1, when to skip it, and how to evaluate reasoning.","href":"/blog/chain-of-thought-prompting-ai-2025","cat":"Blog"},{"title":"Production LLM Monitoring Checklist for 2026: 10 Items Before You Ship","desc":"10-item production LLM monitoring checklist for 2026: OTel, eval gates, drift alerts, PII redaction, A/B rollback, runbooks. Vendor-neutral.","href":"/blog/production-llm-monitoring-checklist-2026","cat":"Blog"},{"title":"Best MCP Gateways in 2026: 6 Production Layers + Companion Instrumentation","desc":"Cloudflare MCP, Bifrost (Maxim), Composio, Smithery, MCP Inspector CLI, and Agent Command Center compared on registration, observability, auth, and OTel.","href":"/blog/best-mcp-gateways-2026","cat":"Blog"},{"title":"Linking Prompt Management with Tracing in 2026: Closing the Loop","desc":"Linking prompt management with tracing in 2026: OTel attribute model, version pinning, A/B variant tags, drift attribution, and eval replay patterns.","href":"/blog/link-prompt-management-tracing-2026","cat":"Blog"},{"title":"AI Explainability in 2026: Tools, Techniques, and Frameworks","desc":"AI explainability in 2026: SHAP, LIME, attention maps, chain-of-thought audits, mechanistic interpretability, and tools that satisfy EU AI Act.","href":"/blog/ai-explainability-tools-techniques-2025","cat":"Blog"},{"title":"Coefficient of Determination (R²) in 2026: How to Interpret It","desc":"How to interpret R² in regression in 2026: when 0.4 is great, when 0.9 means overfitting, the negative-R² trap, and the four metrics you must pair with it.","href":"/blog/coefficient-of-determination-what-it-tells-us-about-our-model","cat":"Blog"},{"title":"Best Text-to-Image AI Models in 2026: 7 Tools Compared","desc":"Compare the 7 best text-to-image AI models in 2026. GPT-image-1, Midjourney v7, FLUX.1, Imagen 4, Stable Diffusion 3.5, Ideogram 3.0, Recraft V3.","href":"/blog/text-to-photo-llm-2025","cat":"Blog"},{"title":"Trace and Debug Multi-Agent Systems in 2026: Production Guide","desc":"Trace, debug, evaluate multi-agent AI systems 2026 with traceAI, OpenTelemetry spans, rubric scoring. Code, span tree, three real failure cases.","href":"/blog/trace-debug-multi-agent-systems-observability-guide","cat":"Blog"},{"title":"AWS Bedrock in 2026: Models, Agents, Guardrails, and Evaluation","desc":"AWS Bedrock in 2026 guide. Claude on Bedrock, Titan, Llama 4, Mistral, Cohere, AI21, Bedrock Agents, Knowledge Bases, Guardrails, plus eval and tracing.","href":"/blog/aws-bedrock-the-future-of-ai-development-on-aws","cat":"Blog"},{"title":"F1 Score in 2026: Formula, Variants, When to Use, Sklearn Code","desc":"F1 Score for classification in 2026: harmonic mean of precision and recall, the math, macro vs micro vs weighted, when to use, sklearn example.","href":"/blog/f1-score-evaluating-classifiers-2025","cat":"Blog"},{"title":"What Are Embeddings in LLMs? Complete 2026 Guide","desc":"How embeddings work in LLMs in 2026. Dense vs sparse, training, dimensionality, semantic vs syntactic, where embeddings sit in modern RAG and agent stacks.","href":"/blog/embeddings-llms-2025","cat":"Blog"},{"title":"How to Get an OpenAI API Key in 2026: 5-Minute Setup","desc":"Generate an OpenAI API key in 2026 with GPT-5 access. Step-by-step setup, secure storage, billing limits, curl + Python examples, and eval add-ons.","href":"/blog/openai-api-key-2025","cat":"Blog"},{"title":"Synthetic Data for AI in 2026: Guide + Best Tools","desc":"How synthetic data works in 2026: rule based, LLM generated, simulation. Use cases, validation, and the tools that ship the highest quality datasets.","href":"/blog/synthetic-data-guide","cat":"Blog"},{"title":"Human vs LLM Annotation in 2026: Accuracy, Cost, Hybrid","desc":"Human vs LLM annotation in 2026: accuracy, Cohen's kappa, cost per label, scalability, and the hybrid LLM-as-judge workflow that production teams now use.","href":"/blog/human-vs-llm-annotation-2025","cat":"Blog"},{"title":"Visual Language Models in 2026: GPT-5o, Claude Opus 4.7, Gemini 3 Pro","desc":"Visual Language Models 2026: GPT-5o vision, Claude Opus 4.7, Gemini 3 Pro, LLaVA, CLIP, BLIP compared, plus how to evaluate multimodal LLMs in production.","href":"/blog/visual-language-models-2025","cat":"Blog"},{"title":"What Is LlamaIndex? 2026 Guide to Workflows and Deployment","desc":"What LlamaIndex is, how it compares to LangChain, and how to build, deploy and evaluate a Workflow in 2026 with llama-agents, llamactl and traceAI.","href":"/blog/exploring-llamaindex-a-powerful-tool-for-llms","cat":"Blog"},{"title":"Single-Turn vs Multi-Turn Evaluation in 2026: A Practical Split","desc":"When to use single-turn LLM eval vs multi-turn, what each measures, and which OSS and commercial tools support each in 2026 production stacks.","href":"/blog/single-turn-vs-multi-turn-evaluation-2026","cat":"Blog"},{"title":"What is the Microsoft Agent Framework? AutoGen + Semantic Kernel for 2026","desc":"Microsoft Agent Framework is the unified successor to AutoGen and Semantic Kernel for production multi-agent systems on Azure. What it is, how to use.","href":"/blog/what-is-microsoft-agent-framework-2026","cat":"Blog"},{"title":"Model Drift vs Data Drift in 2026: Detection & Mitigation Guide","desc":"Model drift vs data drift in 2026: PSI, KS test, embedding cosine, 7 tools ranked. Detect distribution shift in LLM and ML pipelines early.","href":"/blog/model-vs-data-drift-how-to-identify-and-handle-it","cat":"Blog"},{"title":"Agent Architecture Patterns in 2026: The Five Named Shapes","desc":"Five agent architecture patterns in 2026: ReAct, plan-then-execute, supervisor and workers, graph, event-driven. Four-axis tradeoff per pattern.","href":"/blog/agent-architecture-patterns-2026","cat":"Blog"},{"title":"Data Annotation and Synthetic Data in 2026: The Honest Guide","desc":"Data annotation meets synthetic data in 2026: GANs, VAEs, LLM annotators, self-supervision, RLHF, tooling and pitfalls. With FAGI Annotate & Synthesize.","href":"/blog/data-annotation-synthetic-data-2025","cat":"Blog"},{"title":"Time Series Data Analysis in 2026: Models, Frameworks, Code","desc":"Time series data analysis in 2026: Prophet, Darts, statsforecast, neuralforecast, TimesFM, Chronos. Code, benchmarks, when to use each model.","href":"/blog/time-series-data-analysis-2025","cat":"Blog"},{"title":"Best LLM Feedback Collection Tools in 2026: 6 Compared","desc":"Best LLM feedback collection tools in 2026, judged on the closed loop from thumbs to evaluator calibration to CI gate. 6 platforms compared, FAGI included.","href":"/blog/best-llm-feedback-collection-tools-2026","cat":"Blog"},{"title":"Purpose-Built vs General AI Observability in 2026: Where Each Wins","desc":"Datadog and APM vs Phoenix, Langfuse, FutureAGI. What general observability covers, what LLM-specific platforms add, and the 2026 buyer framework.","href":"/blog/purpose-built-vs-general-ai-observability-2026","cat":"Blog"},{"title":"Best LLM Eval Libraries in 2026: 5 OSS Picks Compared","desc":"Best LLM eval libraries in 2026: rubric registries (FAGI ai-evaluation, DeepEval, Ragas) vs test runners (Phoenix Evals, OpenAI Evals) compared.","href":"/blog/best-llm-eval-libraries-2026","cat":"Blog"},{"title":"Error Analysis for LLM Applications: 2026 Workflow Guide","desc":"A 2026 error analysis workflow for LLM apps. Cluster failure cases, label root causes, prioritize fixes. Concrete dataset, code, and rubrics that ship.","href":"/blog/error-analysis-llm-applications-2026","cat":"Blog"},{"title":"RAG Architecture 2026: Patterns, Code, and Eval","desc":"RAG architecture 2026: agentic RAG, multi-hop, query rewriting, hybrid search, reranking, graph RAG. Real code, Context Adherence and Groundedness eval.","href":"/blog/rag-architecture-llm-2025","cat":"Blog"},{"title":"AI Model Testing in 2026: A Practical Multi-Model Comparison Guide","desc":"AI model testing in 2026: how to compare LLMs side by side, score quality, catch bias, pick the right model. Workflow, metrics, FAGI Experiment Feature.","href":"/blog/ai-model-testing-2025","cat":"Blog"},{"title":"Evaluating Causality in AI Models in 2026: Methods and Tools","desc":"Evaluating causality in AI models in 2026. Counterfactuals, RCTs, causal inference for ML, DoWhy, CausalNex, Tetrad, plus LLM causal reasoning eval.","href":"/blog/evaluating-causality-in-ai-models","cat":"Blog"},{"title":"LLM-as-a-Judge in 2026: How It Works, When It Fails","desc":"LLM-as-a-judge in 2026: G-Eval, pairwise, rubric, Cohen's kappa calibration, bias controls, plus tools (FutureAGI, DeepEval, Ragas, Phoenix) compared.","href":"/blog/llm-as-a-judge","cat":"Blog"},{"title":"LangSmith Alternatives in 2026: 6 Honest Picks Compared","desc":"LangSmith alternatives in 2026 compared on cost at scale, LangChain coupling, missing eval, guardrail, and gateway layers. Six honest picks with pricing.","href":"/blog/langsmith-alternatives-2026","cat":"Blog"},{"title":"Stimulus Prompts in 2026: Advanced Prompt Engineering Guide","desc":"Master stimulus prompts in 2026: leading prompts, chain-stimulus, conditioning, prompt chaining, and CI-gated optimization with Future AGI Prompt Optimize.","href":"/blog/stimulus-prompt-guide","cat":"Blog"},{"title":"Synthetic Data Generator in 2026: How It Works, Why You Need One","desc":"What a synthetic data generator does in 2026, the three generation methods, five industry use cases, and how to pick the right tool (with FAGI examples).","href":"/blog/synthetic-data-generator","cat":"Blog"},{"title":"Prompt Caching in 2026: How It Works, Pricing, Wins","desc":"How prompt caching works in 2026 on Anthropic, OpenAI, Gemini, and DeepSeek. Pricing, latency wins on prefix heavy prompts, gotchas, and observability.","href":"/blog/understanding-prompt-caching-for-faster-ai-responses","cat":"Blog"},{"title":"Agent CLI Developer Experience 2026: The 3-Axis DX Test","desc":"Terminal AI coding agents win on three DX axes: plan visibility, tool transparency, rollback discipline. 2026 test for Claude Code, Codex, Aider, Cline.","href":"/blog/agent-cli-developer-experience-2026","cat":"Blog"},{"title":"Model and Prompt Selection in 2026: A Practical Guide","desc":"Pick the right LLM and prompt in 2026: scoring rubric, GPT-5 vs Claude 4.7 vs Gemini 3 trade-offs, automated optimization, and a CI-gated workflow.","href":"/blog/mastering-model-and-prompt-selection-2025","cat":"Blog"},{"title":"What is LLM Monitoring? Alerts, SLOs, Dashboards in 2026","desc":"LLM monitoring is the alerting and dashboard layer on top of observability. Latency, cost, eval pass-rate, drift, and anomaly alerts in 2026.","href":"/blog/what-is-llm-monitoring-2026","cat":"Blog"},{"title":"Benchmarking LLMs for Business Applications in 2026: The Methodology","desc":"How to benchmark LLMs for business in 2026: real-world methodology, metrics that matter beyond MMLU, the modern benchmark stack, 5-step playbook.","href":"/blog/benchmarking-llms-business-applications-2025","cat":"Blog"},{"title":"Non-Deterministic LLM Prompts in 2026: A Practical Guide","desc":"Why LLMs return different answers to the same prompt in 2026, how temperature and top-p actually work, and the four reproducibility levers that matter.","href":"/blog/non-deterministic-llm-prompts-2025","cat":"Blog"},{"title":"Multimodal LLM Tracing in 2026: Image, Audio, Text","desc":"Tracing image, audio, and text spans across multimodal LLM apps in 2026. OTel schema, payload handling, redaction, sampling, and tools that ingest them.","href":"/blog/multimodal-llm-tracing-2026","cat":"Blog"},{"title":"What is LLM Observability? Definition, Stack, OTel in 2026","desc":"LLM observability is traces, OTel GenAI conventions, span-attached evals, cost tracking, and agent graphs. What it is and how to implement it in 2026.","href":"/blog/what-is-llm-observability","cat":"Blog"},{"title":"Dify vs Flowise vs Langflow 2026: 3 No-Code LLM Builders Compared","desc":"Dify, Flowise, and Langflow compared head to head in 2026: license, deployment, RAG depth, agent support, and production readiness.","href":"/blog/dify-vs-flowise-vs-langflow-2026","cat":"Blog"},{"title":"What is LLM Drift? Prompt, Model, and Eval-Score Drift in 2026","desc":"LLM drift is prompt drift, model drift, and eval-score drift in 2026. What it is, how to detect each kind, which tools handle drift on production traces.","href":"/blog/what-is-llm-drift-2026","cat":"Blog"},{"title":"Pipecat Alternatives in 2026: 5 Voice AI Frameworks Compared","desc":"LiveKit Agents, Vapi, Retell, OpenAI Realtime API, and FutureAGI as Pipecat alternatives in 2026. Pricing, OSS license, and real tradeoffs.","href":"/blog/pipecat-alternatives-2026","cat":"Blog"},{"title":"Synthetic Data for LLM Fine-Tuning in 2026: Methods & Stack","desc":"Generate synthetic data to fine-tune LLMs in 2026. Self-Instruct, Constitutional AI, DPO/IPO traces, function calling, and how to evaluate dataset quality.","href":"/blog/synthetic-data-fine-tuning-llms","cat":"Blog"},{"title":"Synthetic Datasets for RAG in 2026: Methods, QA, and Tools","desc":"Synthetic datasets for RAG in 2026: 5 generation methods, quality gates, evaluation metrics, and the 6 tools to use. Includes FutureAGI Dataset workflow.","href":"/blog/synthetic-datasets-rag-2025","cat":"Blog"},{"title":"LLM Hallucination 2026: Causes, Types, and How to Stop It","desc":"What LLM hallucination is in 2026, the six types, why models fabricate, and how to detect each with faithfulness, groundedness, context-adherence scores.","href":"/blog/understanding-llm-hallucination-2025","cat":"Blog"},{"title":"Best Vector Databases for RAG in 2026: 7 Stores Compared","desc":"Pinecone, Milvus, Weaviate, Qdrant, pgvector, Chroma, Vespa for RAG in 2026. Compared on recall, latency, hybrid search, OSS license, eval-fit.","href":"/blog/best-vector-databases-for-rag-2026","cat":"Blog"},{"title":"Best Voice AI Frameworks 2026: 6 Platforms Ranked for Production","desc":"LiveKit Agents, Pipecat, Vapi, Retell, Daily Bots, and OpenAI Realtime API ranked for 2026 by latency, telephony, OSS, and production readiness.","href":"/blog/best-voice-ai-frameworks-2026","cat":"Blog"},{"title":"LiveKit Alternatives in 2026: 5 Voice AI Frameworks Compared","desc":"Pipecat, Vapi, Retell, Daily Bots, and FutureAGI as LiveKit Agents alternatives in 2026. Pricing, OSS license, latency, and real tradeoffs.","href":"/blog/livekit-alternatives-2026","cat":"Blog"},{"title":"Best Embedding Models 2026: NV-Embed, BGE, E5 & OpenAI Compared","desc":"The best embedding models in 2026: NV-Embed-v2, BGE-M3, E5-mistral, OpenAI v3, Voyage 3, Cohere Embed-3. MTEB benchmarks, pricing, and how to pick.","href":"/blog/best-embedding-models-2025","cat":"Blog"},{"title":"LiteLLM vs Alternatives in 2026: Gateway and Proxy Compared","desc":"LiteLLM in 2026 vs Future AGI ACC, Portkey, Helicone, Cloudflare AI Gateway, OpenRouter, vLLM, Ollama: features, security, pick-by-use-case guide.","href":"/blog/litellm-llms-comparison-2025","cat":"Blog"},{"title":"SLM vs LLM in 2026: Cost, Latency, and Quality Compared","desc":"SLM vs LLM in 2026: Gemma 4, Qwen3.5 Small, Phi-4 vs Claude Opus 5, GPT-5.6, Gemini 3.6. Cost, latency, and when to route between them.","href":"/blog/comparison-slm-llm-language-models","cat":"Blog"},{"title":"AI Hallucinations in 2026: Causes, Detection, Prevention","desc":"How AI hallucinations happen in 2026, how to detect them with evaluators, and how RAG, structured output, and guardrails prevent them in production.","href":"/blog/understanding-ai-hallucinations","cat":"Blog"},{"title":"Getting Started with AI Agent Evaluation in 2026 (Future AGI Tutorial)","desc":"Evaluate AI agents in 2026 with Future AGI: fi.evals quickstart, fi.simulate scenarios, traceAI instrumentation, metrics, and production pipeline.","href":"/blog/getting-started-with-agent-evaluation","cat":"Blog"},{"title":"LLM Function Calling in 2026: OpenAI, Anthropic, Structured Outputs","desc":"How LLM function calling works in 2026. JSON Schema, OpenAI tools, Anthropic tools, structured outputs, parallel tool calls, function-call eval.","href":"/blog/llm-function-calling-2025","cat":"Blog"},{"title":"How to Build LLM Agents in 2026: A Production Guide","desc":"Build production LLM agents in 2026: task scoping, model selection (gpt-5, claude-opus-4.5), tools, evals, observability, the orchestration plus eval loop.","href":"/blog/build-llm-agents","cat":"Blog"},{"title":"Building LLMs in Production 2026: A Step-by-Step Playbook","desc":"How to ship LLMs to production in 2026. Covers data, model selection, gpt-5, claude-opus-4-7, eval, observability, scaling, and the FAGI deployment loop.","href":"/blog/building-llms-production-2025","cat":"Blog"},{"title":"Best Free AI Search Engines 2026: Top 7 Ranked & Compared","desc":"Ranked: 7 best free AI search engines for August 2026. Perplexity, ChatGPT Search, Google AI Mode, Brave AI, Andi, Duck.ai and You.com compared.","href":"/blog/free-ai-search-engines","cat":"Blog"},{"title":"AI Agent Evaluation in 2026: Tool Trajectory, Persona Sim, Real Code","desc":"Evaluate AI agents 2026 with task completion, tool trajectory, response quality, multi-turn checks, persona simulation. Real fi.evals + fi.simulate code.","href":"/blog/mastering-evaluation-ai-agents-2025","cat":"Blog"},{"title":"Best Free AI Search Tools in 2026: Perplexity, Phind, You.com","desc":"The 6 best free AI search tools in 2026: Perplexity Free, ChatGPT Search free, Phind, You.com, Brave Search AI, DuckDuckGo AI. Real limits, real strengths.","href":"/blog/ai-search-free-2025","cat":"Blog"},{"title":"AI Search Engines in 2026: Perplexity, You.com, Phind, Kagi, ChatGPT","desc":"The AI search engines that work in 2026 with their free tiers. Compare Perplexity, You.com, Phind, Kagi, ChatGPT Search, Gemini, and Claude web search.","href":"/blog/free-easiest-ai-search-engine","cat":"Blog"},{"title":"LLM Fine-Tuning Techniques 2026: LoRA, QLoRA, SFT, DPO","desc":"LLM fine-tuning techniques in 2026: feature-based, full fine-tune, LoRA, QLoRA, BitFit, SFT, DPO, RLHF, multi-task. When to use each and how to evaluate.","href":"/blog/llm-fine-tuning-techniques-i-ii","cat":"Blog"},{"title":"Mean Squared Error (MSE) in Machine Learning: Formula, RMSE, MAE, R-Squared","desc":"Complete MSE guide for 2026. Formula, Python example, when MSE beats MAE or RMSE, R-squared comparison, outlier sensitivity, neural network loss use cases.","href":"/blog/mean-squared-error-2025","cat":"Blog"},{"title":"Hard Prompt vs Soft Prompt in 2026: Differences and When to Use","desc":"Hard prompts vs soft prompts in 2026: prompt tuning, prefix tuning, P-tuning, LoRA for prompts. Decision guide, code, and benchmarks for production teams.","href":"/blog/hard-prompt-vs-soft-prompt-2025","cat":"Blog"},{"title":"What is Eval-Driven Development? The TDD-for-LLMs Workflow in 2026","desc":"Eval-driven development writes the eval first, then iterates the prompt against it. The TDD analog for LLM apps, the cycle, and how teams adopt it in 2026.","href":"/blog/what-is-eval-driven-development-2026","cat":"Blog"},{"title":"AI for Creating Dashboards in 2026: Tools and Workflow","desc":"AI for creating dashboards in 2026: Hex Magic, Mode AI, Power BI Copilot, Tableau Pulse, Looker compared, plus a six-step build workflow and observability.","href":"/blog/ai-for-creating-dashboards","cat":"Blog"},{"title":"K-Nearest Neighbor (KNN) in 2026: How It Works and When to Use It","desc":"Learn how K-Nearest Neighbor (KNN) works in 2026. Distance metrics, parameter tuning, and when to use KNN vs decision trees, SVMs, and neural networks.","href":"/blog/k-nearest-neighbor","cat":"Blog"},{"title":"RAG Prompting to Reduce Hallucination: 6 Techniques 2026","desc":"Six RAG prompting patterns that reduce hallucination, with example prompts, retrieval grounding, and Context Adherence + Groundedness eval code.","href":"/blog/rag-prompting-to-reduce-hallucination","cat":"Blog"},{"title":"Advanced RAG Chunking Techniques in 2026: Late, Semantic, and Parent-Child","desc":"Ranked RAG chunking strategies for 2026: late chunking, semantic, hierarchical, parent-child, sliding window. Code, tradeoffs, how to evaluate retrieval.","href":"/blog/advanced-chunking-techniques-for-rag","cat":"Blog"},{"title":"Agentic AI Workflows in 2026: Architecture, Reliability, Use Cases","desc":"Agentic AI workflows in 2026: 4 architecture patterns, 6 reliability metrics, use cases in healthcare, finance, ops with traceable, evaluable agents.","href":"/blog/agentic-ai-workflows-game-changer-automation-ethics-future","cat":"Blog"},{"title":"Fine-Tune Prompts (Not Models) for LLMs in 2026: Full Guide","desc":"Fine-tune prompts (not weights) to lift LLM accuracy in 2026. Covers DSPy, prompt-opt loops, FAGI Prompt-Opt, MIPRO, and a runnable eval loop you can ship.","href":"/blog/fine-tune-prompts-llm-2025","cat":"Blog"},{"title":"LLM vs GPT 2026: Key Differences, How They Work, and When to Use Each","desc":"LLM vs GPT in 2026 explained: GPT is one family of LLM. Definitions, architecture, GPT-5.6 vs Claude Opus 5 vs Gemini 3.x, and how to evaluate any model.","href":"/blog/llm-vs-gpt","cat":"Blog"},{"title":"R-Squared (R²) Explained: Formula, Interpretation, Pitfalls","desc":"What R-squared means, how to compute it, when adjusted R² helps, when to switch to RMSE/MAE, and why LLM evaluation needs different metrics.","href":"/blog/r-squared-model-accuracy-2025","cat":"Blog"},{"title":"Intelligent AI Agents in 2026: How They Work and 6 Use Cases","desc":"What intelligent agents are in 2026: architecture, RL foundations, multi-agent systems, evaluation, observability, 5 production use cases.","href":"/blog/exploring-intelligent-agents-ai-automation-decision-making","cat":"Blog"},{"title":"Top Open-Source LLMs in 2026: Llama 4, DeepSeek R2, Qwen 3","desc":"The 7 leading open-source LLMs in 2026: Llama 4, DeepSeek R2, Qwen 3, Mistral, Phi-5, Gemma 3, OLMo. Licenses, hardware, benchmarks, and how to choose.","href":"/blog/top-open-source-llms-in-2025-driving-innovation-in-ai","cat":"Blog"},{"title":"Continued LLM Pretraining in 2026: Frameworks, Strategies, Evaluation","desc":"Continued LLM pretraining in 2026: Megatron-LM, DeepSpeed, Axolotl, NeMo, Unsloth. Domain adaptation, catastrophic forgetting, evaluation with Future AGI.","href":"/blog/continued-llm-pretraining","cat":"Blog"},{"title":"Productionize Agentic Apps in 2026: 9-Step Playbook","desc":"Ship agentic apps to production in 2026: orchestration, eval gates, traceAI observability, guardrails, MCP, and rollback. 9 steps with code and metrics.","href":"/blog/how-to-productionize-agentic-applications","cat":"Blog"},{"title":"No-Code LLM AI in 2026: Platforms, Patterns, and Buyer Guide","desc":"How no-code LLM AI works in 2026, the platforms that ship, what to look for, and how to evaluate the AI you build. Citizen developer's pragmatic guide.","href":"/blog/no-code-llm-ai","cat":"Blog"},{"title":"RAG and Perplexity in 2026: Metric vs. Product, Plus What to Use","desc":"Perplexity for RAG in 2026: the metric vs Perplexity.ai the product. When perplexity is the right LLM score, when faithfulness wins, plus the eval stack.","href":"/blog/rag-llm-perplexity-2025","cat":"Blog"},{"title":"Small Language Models for Agentic AI in 2026: SLM Lineup + Build Guide","desc":"The 2026 SLM lineup for agentic AI (Phi-4, Llama 3.2, Ministral, Gemma 2, Qwen 2.5) plus a build pattern for modular multi-agent workflows.","href":"/blog/small-language-models-agentic-ai-2025","cat":"Blog"},{"title":"Prompt Engineering Careers 2026: Roles, Salaries, Skills","desc":"Prompt engineering careers in 2026: actual job titles, illustrative salary ranges, the eight skills hiring managers test, and where to start.","href":"/blog/agi-careers-prompt-engineering-opportunities","cat":"Blog"},{"title":"Future Trends in Generative AI for 2026: 7 Shifts to Track","desc":"Seven generative AI trends to track in 2026: agentic workflows, multimodal, custom evals, MCP, on-device, routing, and closed-loop eval with traceAI.","href":"/blog/future-trends-generative-ai-2025","cat":"Blog"},{"title":"GenAI Plus No-Code Platforms in 2026: A Buyer's Guide","desc":"How generative AI and no-code platforms combine in 2026: GPT-5, Claude 4.7, Gemini 3 inside Dify, Flowise, Langflow, n8n, Vapi. Ship vs avoid.","href":"/blog/generative-ai-no-code-platforms-empowering-creativity-innovation","cat":"Blog"},{"title":"RAG vs Fine-Tuning in 2026: Which AI Strategy Should You Pick?","desc":"RAG vs fine-tuning in 2026: decision matrix on data freshness, cost, latency, accuracy, governance, and how to evaluate either path with Future AGI.","href":"/blog/rag-vs-fine-tuning-which-ai-training-strategy-is-right","cat":"Blog"},{"title":"Real-Time Learning in LLMs (2026): Online Learning Methods Explained","desc":"How real-time and online learning works in LLMs in 2026: continual learning, RLHF, DPO, GRPO, LoRA, MoE, retrieval-augmented adaptation, and trade-offs.","href":"/blog/real-time-learning-in-large-language-models-llms","cat":"Blog"},{"title":"User Feedback Loops in 2026: Closing the AI Data Improvement Cycle","desc":"Integrate user feedback into automated data layers in 2026. Five steps: capture, classify, prioritize, augment datasets, gate releases on regression.","href":"/blog/integrating-user-feedback-automated-data-layers","cat":"Blog"},{"title":"AI Agents in 2026: The Good, the Bad, and the Unknown","desc":"What 2026 AI agents do well, where they still fail, and the open questions. A grounded read for teams shipping autonomous LLM systems.","href":"/blog/ai-agents-the-good-the-bad-and-the-unknown","cat":"Blog"},{"title":"Prompt Engineering in 2026: 10 Patterns That Actually Work","desc":"Prompt engineering patterns that actually move LLM performance in 2026: CoT, ToT, structured outputs, XML tags, multi-shot, plus tools and benchmarks.","href":"/blog/effective-prompt-engineering-maximize-llm-performance","cat":"Blog"},{"title":"Dynamic Prompts in 2026: Template, Variables, and Runtime Context","desc":"Dynamic prompts in 2026: template engines, variable injection, runtime context, versioning, and evaluation. With code, failure modes, and an eval harness.","href":"/blog/dynamic-prompts","cat":"Blog"},{"title":"Fine-Tuning LLMs in 2026: LoRA, QLoRA, DPO, GRPO Compared","desc":"2026 guide to fine-tuning LLMs: LoRA vs QLoRA, DPO vs RLHF vs GRPO, and when to fine-tune open-weight models instead of prompting alone.","href":"/blog/fine-tuning-llms-unlocking-peak-performance","cat":"Blog"},{"title":"How to Evaluate LLMs in 2026: Metrics, Frameworks, Pipelines","desc":"How to evaluate LLMs in 2026. Pick use-case metrics, score with judges + heuristics, gate CI, and run continuous production evals in under 200 lines.","href":"/blog/how-to-evaluate-large-language-models-llms","cat":"Blog"},{"title":"Best Books for Learning LLM Training in 2026: 10 Picks","desc":"Best books and free courses to learn LLM training in 2026: Sutton's RL, Goodfellow Deep Learning, Jurafsky SLP, Karpathy CS25, playbooks.","href":"/blog/large-language-model-training-books-2025","cat":"Blog"},{"title":"Automated Error Detection for Generative AI in 2026","desc":"Automated error detection for generative AI in 2026. Compares the top platforms, real traceAI + fi.evals patterns, and rollout playbook.","href":"/blog/leveraging-automated-error-detection-in-generative-ai-workflows","cat":"Blog"},{"title":"Best Open-Weight LLMs 2026: Llama 4, DeepSeek R2, Qwen 3 Compared","desc":"Compare the top open-weight LLMs in 2026: Llama 4.x, DeepSeek R2, Qwen 3, Mistral, Phi family. Benchmarks, licensing, hardware, and how to evaluate yours.","href":"/blog/open-source-llms-2025","cat":"Blog"},{"title":"LLM Experimentation in 2026: Best Practices and Tools","desc":"LLM experimentation 2026: 6 best practices, 5 trends (LoRA, multimodal, MoE), a ranked stack for prompt-opt, evals, tracing. Production-ready guide.","href":"/blog/optimizing-llm-experimentation-best-practices","cat":"Blog"},{"title":"Real-Time LLM Performance Monitoring in 2026: 7 Tools Ranked","desc":"Real-time LLM monitoring in 2026. FutureAGI, Langfuse, Phoenix, Helicone, OpenLIT, Datadog, and New Relic ranked on latency, eval depth, and OTel support.","href":"/blog/real-time-monitoring-of-llm-performance","cat":"Blog"},{"title":"What Is Prompt Tuning? 2026 Guide vs Prompt Engineering","desc":"Prompt tuning explained for 2026. Soft prompts, P-Tuning, prefix tuning, plus how it differs from prompt engineering and fine-tuning on gpt-5 and Llama 4.","href":"/blog/what-is-prompt-tuning","cat":"Blog"},{"title":"Automate LLM Data Annotation in 2026: A Practical Guide","desc":"How to automate LLM data annotation in 2026. Calibrated LLM judges, compound vs single calls, gold-set bootstrapping, Future AGI synthetic tooling.","href":"/blog/automating-data-annotation-for-llms","cat":"Blog"},{"title":"Contextual Chatbots 2026: Customer Engagement at Scale","desc":"Build contextual chatbots 2026: NLP, ML, RAG, evaluation, observability. Top tools compared, FAGI eval stack, real-time guardrails for production.","href":"/blog/contextual-chatbots-customer-engagement","cat":"Blog"},{"title":"Self-Learning Agents in 2026: Build a Self-Improving Agent Loop with FAGI","desc":"Self-learning AI agents in 2026: build the eval-and-optimize loop with Future AGI fi.opt optimizers, fi.evals scoring, and traceAI tracing in production.","href":"/blog/self-learning-agents-ai-transformation-futureagi","cat":"Blog"},{"title":"How to Reduce LLM Hallucinations in 2026: 7 Proven Strategies","desc":"Reduce LLM hallucinations in 2026 with seven proven strategies: RAG grounding, uncertainty estimation, fine tuning, adversarial training, live eval.","href":"/blog/taming-hallucination-beast-strategies-reliable-llms","cat":"Blog"},{"title":"RAG Summarization 2026: Patterns, Code + Long-Context Tradeoffs","desc":"RAG summarization in 2026: stuff, map-reduce, refine, RAPTOR, GraphRAG. Long-context vs RAG decision matrix with thresholds plus faithfulness eval code.","href":"/blog/rag-summarization","cat":"Blog"},{"title":"Scaling High-Fidelity Synthetic Data Generation with Future AGI","desc":"Multi-agent framework for generating high-quality, diverse, privacy-preserving synthetic datasets, with perfect quality scores on standard benchmarks.","href":"/research/synthetic-data-generation","cat":"Research"},{"title":"AgentCompass: Towards Reliable Evaluation of Agentic Workflows in Production","desc":"AgentCompass: first evaluation framework purpose-built for monitoring and debugging agentic workflows in production, SOTA on the TRAIL benchmark.","href":"/research/agent-compass","cat":"Research"},{"title":"Protect: Towards Robust Guardrailing Stack for Trustworthy Enterprise LLM Systems","desc":"Protect: a natively multi-modal guardrailing model across text, image, and audio with SOTA results on toxicity, sexism, privacy, and prompt injection.","href":"/research/protect-guardrailing-stack","cat":"Research"},{"title":"FutureAGI's Evaluation Framework: Precision, Adaptability, and Explainability","desc":"Multi-agent evaluation system with SOTA performance across NLI, commonsense reasoning, toxicity classification, and vision-language tasks.","href":"/research/multimodal-evaluation-framework","cat":"Research"},{"title":"60% fewer chatbot hallucinations with AI observability","desc":"A leading SaaS provider used Trace AI to cut factual inaccuracies by 60% and reduce LLM API costs by 22% in their customer support chatbot.","href":"/customers/ai-observability-customer-support","cat":"Case Study"},{"title":"Benchmarking LLMs for customer support in 3 days","desc":"How Future AGI's observability platform helped benchmark Mistral, Claude, and GPT-4o across 12+ metrics in just 3 days.","href":"/customers/benchmarking-llms-customer-support","cat":"Case Study"},{"title":"Autonomous agents in production: 95% task completion rate","desc":"An AI automation company used Future AGI to test multi-step workflows, detect loops, and achieve 95% task completion in production.","href":"/customers/autonomous-agent-workflows","cat":"Case Study"},{"title":"25% higher response rates with intelligent prompt evaluation","desc":"An AI SDR company used Future AGI to optimize lead generation prompts, achieving 25% better response rates and 10x evaluation scale.","href":"/customers/ai-sdr-lead-generation","cat":"Case Study"},{"title":"10x faster quiz validation for EdTech at scale","desc":"An EdTech company used Future AGI to validate AI-generated quizzes 10x faster with 80% less manual effort.","href":"/customers/edtech-quiz-validation","cat":"Case Study"},{"title":"Building accurate fintech chatbots with evaluation & observability","desc":"A fintech platform used Future AGI to reduce chatbot errors by 25% and boost first-contact resolution by 15% for financial queries.","href":"/customers/fintech-chatbot-accuracy","cat":"Case Study"},{"title":"Computer-use agents that click with 99% accuracy","desc":"An enterprise used Future AGI to simulate UI workflows, block destructive actions, and achieve 99% click accuracy for CUA agents.","href":"/customers/computer-use-agent-testing","cat":"Case Study"},{"title":"Building a coding agent that ships safe code to production","desc":"A startup building an AI coding agent used Future AGI to evaluate generated code across languages, block destructive ops, and ship reliably.","href":"/customers/coding-agent-safety","cat":"Case Study"},{"title":"10x HR productivity with AI-powered knowledge optimization","desc":"Future AGI helped an enterprise HR team achieve 65% faster document creation and 99% compliance through intelligent evaluation.","href":"/customers/hr-productivity-optimization","cat":"Case Study"},{"title":"90% less manual effort in meeting summarization evaluation","desc":"How Future AGI's evaluation framework automated model selection for meeting summarization with objective, scalable metrics.","href":"/customers/meeting-summarization-evaluation","cat":"Case Study"},{"title":"10x faster image AI evaluation for creative workflows","desc":"An AI comic generation company used Future AGI to automate image evaluation, achieving 85% less manual effort and 10x throughput.","href":"/customers/optimizing-image-ai","cat":"Case Study"},{"title":"Elevating SQL accuracy for retail analytics at scale","desc":"A Fortune-50 retailer used Future AGI to validate SQL agents, achieving 10x faster query validation and 90% fewer errors.","href":"/customers/sql-query-validation-retail","cat":"Case Study"},{"title":"Reinventing Tier-1 support with Prompt Playground","desc":"A global tech leader cut Tier-1 resolution times by 30% and saved $1.2M annually using Future AGI's Prompt Playground.","href":"/customers/tier-1-customer-support","cat":"Case Study"},{"title":"Voice AI quality at scale: 40% fewer call failures","desc":"A voice AI platform used Future AGI to test diverse personas, evaluate STT/TTS/LLM independently, and cut call failures by 40%.","href":"/customers/voice-agent-quality","cat":"Case Study"},{"title":"Advanced RAG Patterns","desc":"Standard RAG breaks at enterprise scale: missed answers, costs, compliance risks. Architecture patterns to build retrieval systems you can rely on.","href":"/ebooks/advanced-rag-patterns","cat":"eBook"},{"title":"The Agentic RAG Playbook","desc":"Transform RAG theory into product-ready enterprise solutions that deliver measurable business impact across agentic workflows and retrieval systems.","href":"/ebooks/mastering-agentic-rag","cat":"eBook"},{"title":"Mastering AI Agent Evaluation","desc":"AI agents are easy to spin up and dangerously hard to trust in production. A concrete evaluation playbook to turn messy agents into controlled systems.","href":"/ebooks/mastering-ai-agent-evaluation","cat":"eBook"},{"title":"Handbook","desc":"The Flight Manual - how we work, what we believe","href":"/handbook/","cat":"Pages"},{"title":"Contributing","desc":"A guide for external contributors - how to help, what we're looking for, and how we work together.","href":"/handbook/community/contributing","cat":"Handbook"},{"title":"Open Source","desc":"Our relationship with open source - what we contribute, how we think about it, and why it matters.","href":"/handbook/community/open-source","cat":"Handbook"},{"title":"Benefits & Perks","desc":"What we offer beyond salary - health, time off, equipment, and the things that make work sustainable.","href":"/handbook/compensation/benefits","cat":"Handbook"},{"title":"Pay Philosophy","desc":"How we think about compensation - transparency, fairness, and paying well.","href":"/handbook/compensation/pay-philosophy","cat":"Handbook"},{"title":"Communication","desc":"Channels, norms, and expectations for how we communicate.","href":"/handbook/culture/communication","cat":"Handbook"},{"title":"Crew Manifest","desc":"Every mission flies with a team. This is ours.","href":"/handbook/crew/launch-team","cat":"Handbook"},{"title":"Decision Making","desc":"How we make decisions - DRIs, RFCs, and disagree-and-commit.","href":"/handbook/culture/decision-making","cat":"Handbook"},{"title":"Support Philosophy","desc":"How we help customers succeed - fast responses, real engineers, no ticket deflection.","href":"/handbook/customer-success/support-philosophy","cat":"Handbook"},{"title":"Working Principles","desc":"How we work day-to-day - async-first, high autonomy, bias toward shipping.","href":"/handbook/culture/working-principles","cat":"Handbook"},{"title":"Customer Onboarding","desc":"How we bring new customers from signup to production value as fast as possible.","href":"/handbook/customer-success/customer-onboarding","cat":"Handbook"},{"title":"Design Principles","desc":"The principles that guide how we design our product, brand, and experiences.","href":"/handbook/design/design-principles","cat":"Handbook"},{"title":"Brand Identity","desc":"Our visual identity, voice, and the aesthetic principles behind Future AGI.","href":"/handbook/design/brand-identity","cat":"Handbook"},{"title":"Code Review","desc":"Code review philosophy, what to look for, and turnaround expectations.","href":"/handbook/engineering/code-review","cat":"Handbook"},{"title":"Development Process","desc":"How features go from idea to production - branches, PRs, CI/CD, and shipping.","href":"/handbook/engineering/development-process","cat":"Handbook"},{"title":"Tech Stack","desc":"The technologies we use and why we chose them.","href":"/handbook/engineering/tech-stack","cat":"Handbook"},{"title":"AI Safety Philosophy","desc":"How we think about AI safety - engineering rigor, not fear.","href":"/handbook/mission/ai-safety-philosophy","cat":"Handbook"},{"title":"Our Values","desc":"The principles that guide how we work, build, and make decisions.","href":"/handbook/mission/our-values","cat":"Handbook"},{"title":"What We Build","desc":"An overview of the Future AGI platform and its five-stage pipeline.","href":"/handbook/mission/what-we-build","cat":"Handbook"},{"title":"Why We Exist","desc":"The problem we're solving and why it matters now more than ever.","href":"/handbook/mission/why-we-exist","cat":"Handbook"},{"title":"Tools We Use","desc":"The tools and services that power our day-to-day operations.","href":"/handbook/operations/tools-we-use","cat":"Handbook"},{"title":"Developer Experience","desc":"How we think about the end-to-end experience for developers using our platform.","href":"/handbook/growth/developer-experience","cat":"Handbook"},{"title":"How We Prioritize","desc":"How the product team decides what to build next.","href":"/handbook/product/how-we-prioritize","cat":"Handbook"},{"title":"Data Handling","desc":"How we handle customer data - collection, storage, retention, and deletion.","href":"/handbook/security/data-handling","cat":"Handbook"},{"title":"How We Grow","desc":"Our growth philosophy - product-led, developer-first, and honest.","href":"/handbook/growth/how-we-grow","cat":"Handbook"},{"title":"Hiring Process","desc":"End-to-end hiring - how we source, interview, and decide.","href":"/handbook/people/hiring-process","cat":"Handbook"},{"title":"Security Practices","desc":"How we protect our product, our customers' data, and our own infrastructure.","href":"/handbook/security/security-practices","cat":"Handbook"},{"title":"Onboarding","desc":"Your first 30/60/90 days at Future AGI.","href":"/handbook/people/onboarding","cat":"Handbook"}]</script> </div> </div> </div>  <script type="module" src="/_astro/Header.astro_astro_type_script_index_0_lang.CCSGs3CY.js"></script> <main class="startups-page" data-astro-cid-myseyla7> <!-- ===================== HERO: A ROCKET IS BUILDING ===================== --> <section class="su-hero" data-astro-cid-myseyla7> <canvas id="rocket-canvas" class="su-hero-canvas" data-astro-cid-myseyla7></canvas> <div class="su-hero-fade" data-astro-cid-myseyla7></div> <!-- HUD corners --> <div class="su-hud su-hud-tl" data-astro-cid-myseyla7></div> <div class="su-hud su-hud-tr" data-astro-cid-myseyla7></div> <div class="su-hud su-hud-bl" data-astro-cid-myseyla7></div> <div class="su-hud su-hud-br" data-astro-cid-myseyla7></div> <!-- Content overlay --> <div class="su-hero-content" data-astro-cid-myseyla7> <div class="su-hero-badge" data-astro-cid-myseyla7> <span class="su-bracket" data-astro-cid-myseyla7>[</span> <span class="su-badge-text" data-astro-cid-myseyla7>STARTUP PROGRAM</span> <span class="su-badge-ver" data-astro-cid-myseyla7>v2.0</span> <span class="su-bracket" data-astro-cid-myseyla7>]</span> </div> <h1 class="su-hero-title" data-astro-cid-myseyla7> <span class="su-word" style="--i:0" data-astro-cid-myseyla7>We</span> <span class="su-word" style="--i:1" data-astro-cid-myseyla7>back</span> <span class="su-word" style="--i:2" data-astro-cid-myseyla7>builders.</span> </h1> <p class="su-hero-sub" data-astro-cid-myseyla7>
$6K in credits. 6 months Pro. Direct engineering support.<br data-astro-cid-myseyla7>
Everything you need to launch AI that doesn't hallucinate.
</p> <div class="su-hero-cta" data-astro-cid-myseyla7> <button class="su-btn-primary open-apply-modal" data-astro-cid-myseyla7> <span class="su-btn-prompt" data-astro-cid-myseyla7>$</span> initiate_launch --apply
</button> <a href="#launchpad" class="su-btn-ghost" data-astro-cid-myseyla7> <span class="su-btn-prompt" data-astro-cid-myseyla7>&gt;</span> view_manifest
</a> </div> <div class="su-hero-stats" data-astro-cid-myseyla7> <div class="su-stat" data-astro-cid-myseyla7> <span class="su-stat-val" data-astro-cid-myseyla7>500+</span> <span class="su-stat-label" data-astro-cid-myseyla7>startups enrolled</span> </div> <span class="su-stat-sep" data-astro-cid-myseyla7>|</span> <div class="su-stat" data-astro-cid-myseyla7> <span class="su-stat-val" data-astro-cid-myseyla7>48h</span> <span class="su-stat-label" data-astro-cid-myseyla7>approval time</span> </div> <span class="su-stat-sep" data-astro-cid-myseyla7>|</span> <div class="su-stat" data-astro-cid-myseyla7> <span class="su-stat-val" data-astro-cid-myseyla7>$6K</span> <span class="su-stat-label" data-astro-cid-myseyla7>free credits</span> </div> </div> </div> </section> <!-- ===================== THE LAUNCHPAD: PERKS ===================== --> <section id="launchpad" class="su-section su-launchpad" data-astro-cid-myseyla7> <div class="su-container" data-astro-cid-myseyla7> <div class="su-section-header" data-astro-cid-myseyla7> <div class="su-divider-line" data-astro-cid-myseyla7></div> <span class="su-section-code" data-astro-cid-myseyla7>SYS-01</span> <span class="su-section-label" data-astro-cid-myseyla7>THE LAUNCHPAD</span> <div class="su-divider-line" data-astro-cid-myseyla7></div> </div> <h2 class="su-section-title" data-astro-cid-myseyla7>Mission manifest</h2> <p class="su-section-desc" data-astro-cid-myseyla7>What's included in the startup program.</p> <div class="su-perks-grid" data-astro-cid-myseyla7> <div class="su-perk" style="--d:0s" data-astro-cid-myseyla7> <div class="su-perk-header" data-astro-cid-myseyla7> <span class="su-perk-code" data-astro-cid-myseyla7>CRED-6K</span> <span class="su-perk-dot" data-astro-cid-myseyla7></span> <span class="su-perk-status" data-astro-cid-myseyla7>ACTIVE</span> </div> <div class="su-perk-value" data-astro-cid-myseyla7>$6,000</div> <div class="su-perk-label" data-astro-cid-myseyla7>Platform Credits</div> <p class="su-perk-desc" data-astro-cid-myseyla7>One year of free credits to build, test, and iterate without worrying about cost.</p> <div class="su-perk-border" data-astro-cid-myseyla7></div> </div><div class="su-perk" style="--d:0.08s" data-astro-cid-myseyla7> <div class="su-perk-header" data-astro-cid-myseyla7> <span class="su-perk-code" data-astro-cid-myseyla7>PRO-6MO</span> <span class="su-perk-dot" data-astro-cid-myseyla7></span> <span class="su-perk-status" data-astro-cid-myseyla7>ACTIVE</span> </div> <div class="su-perk-value" data-astro-cid-myseyla7>6 Months</div> <div class="su-perk-label" data-astro-cid-myseyla7>Pro Plan Access</div> <p class="su-perk-desc" data-astro-cid-myseyla7>Every feature unlocked: advanced evals, custom guardrails, priority queues.</p> <div class="su-perk-border" data-astro-cid-myseyla7></div> </div><div class="su-perk" style="--d:0.16s" data-astro-cid-myseyla7> <div class="su-perk-header" data-astro-cid-myseyla7> <span class="su-perk-code" data-astro-cid-myseyla7>SLACK-ENG</span> <span class="su-perk-dot" data-astro-cid-myseyla7></span> <span class="su-perk-status" data-astro-cid-myseyla7>ACTIVE</span> </div> <div class="su-perk-value" data-astro-cid-myseyla7>Direct Line</div> <div class="su-perk-label" data-astro-cid-myseyla7>Engineering Slack</div> <p class="su-perk-desc" data-astro-cid-myseyla7>A private channel with our eng team. Ask anything, ship faster.</p> <div class="su-perk-border" data-astro-cid-myseyla7></div> </div><div class="su-perk" style="--d:0.24s" data-astro-cid-myseyla7> <div class="su-perk-header" data-astro-cid-myseyla7> <span class="su-perk-code" data-astro-cid-myseyla7>OFC-HRS</span> <span class="su-perk-dot" data-astro-cid-myseyla7></span> <span class="su-perk-status" data-astro-cid-myseyla7>ACTIVE</span> </div> <div class="su-perk-value" data-astro-cid-myseyla7>Weekly</div> <div class="su-perk-label" data-astro-cid-myseyla7>Founder Office Hours</div> <p class="su-perk-desc" data-astro-cid-myseyla7>Strategy sessions with our founders on AI architecture, evals, and go-to-market.</p> <div class="su-perk-border" data-astro-cid-myseyla7></div> </div> </div> </div> </section> <!-- ===================== FLIGHT CHECKLIST: ELIGIBILITY ===================== --> <section class="su-section su-checklist-section" data-astro-cid-myseyla7> <div class="su-container" data-astro-cid-myseyla7> <div class="su-section-header" data-astro-cid-myseyla7> <div class="su-divider-line" data-astro-cid-myseyla7></div> <span class="su-section-code" data-astro-cid-myseyla7>SYS-02</span> <span class="su-section-label" data-astro-cid-myseyla7>FLIGHT CHECKLIST</span> <div class="su-divider-line" data-astro-cid-myseyla7></div> </div> <h2 class="su-section-title" data-astro-cid-myseyla7>Pre-flight check</h2> <p class="su-section-desc" data-astro-cid-myseyla7>Eligibility requirements for the program.</p> <div class="su-checklist" data-astro-cid-myseyla7> <div class="su-terminal-bar" data-astro-cid-myseyla7> <span class="su-terminal-dot" data-astro-cid-myseyla7></span> <span class="su-terminal-dot" data-astro-cid-myseyla7></span> <span class="su-terminal-dot" data-astro-cid-myseyla7></span> <span class="su-terminal-title" data-astro-cid-myseyla7>eligibility.sh</span> </div> <div class="su-terminal-body" data-astro-cid-myseyla7> <div class="su-terminal-line su-terminal-comment" data-astro-cid-myseyla7># Run pre-flight eligibility checks</div> <div class="su-terminal-line su-terminal-comment" data-astro-cid-myseyla7># All checks must PASS to qualify</div> <div class="su-terminal-line" data-astro-cid-myseyla7>&nbsp;</div> <div class="su-check-item" style="--d:0s" data-astro-cid-myseyla7> <span class="su-check-icon su-check-pass" data-astro-cid-myseyla7> [✓] </span> <span class="su-check-text" data-astro-cid-myseyla7>Less than $10M in total funding raised</span> <span class="su-check-status su-status-pass" data-astro-cid-myseyla7> PASS </span> </div><div class="su-check-item" style="--d:0.1s" data-astro-cid-myseyla7> <span class="su-check-icon su-check-pass" data-astro-cid-myseyla7> [✓] </span> <span class="su-check-text" data-astro-cid-myseyla7>Company is less than 5 years old</span> <span class="su-check-status su-status-pass" data-astro-cid-myseyla7> PASS </span> </div><div class="su-check-item" style="--d:0.2s" data-astro-cid-myseyla7> <span class="su-check-icon su-check-pass" data-astro-cid-myseyla7> [✓] </span> <span class="su-check-text" data-astro-cid-myseyla7>Building a product that uses LLMs or AI agents</span> <span class="su-check-status su-status-pass" data-astro-cid-myseyla7> PASS </span> </div><div class="su-check-item" style="--d:0.30000000000000004s" data-astro-cid-myseyla7> <span class="su-check-icon su-check-pass" data-astro-cid-myseyla7> [✓] </span> <span class="su-check-text" data-astro-cid-myseyla7>Has a live product or working prototype</span> <span class="su-check-status su-status-pass" data-astro-cid-myseyla7> PASS </span> </div><div class="su-check-item" style="--d:0.4s" data-astro-cid-myseyla7> <span class="su-check-icon su-check-note" data-astro-cid-myseyla7> [~] </span> <span class="su-check-text" data-astro-cid-myseyla7>Not sure? Apply anyway. We review every application.</span> <span class="su-check-status su-status-note" data-astro-cid-myseyla7> NOTE </span> </div> <div class="su-terminal-line" data-astro-cid-myseyla7>&nbsp;</div> <div class="su-terminal-line su-terminal-output" data-astro-cid-myseyla7> <span class="su-terminal-prompt" data-astro-cid-myseyla7>$</span> echo "Ready for launch? Apply now."
</div> </div> </div> </div> </section> <!-- ===================== TRANSMISSION LOG: SOCIAL PROOF ===================== --> <section class="su-section su-transmissions" data-astro-cid-myseyla7> <div class="su-container" data-astro-cid-myseyla7> <div class="su-section-header" data-astro-cid-myseyla7> <div class="su-divider-line" data-astro-cid-myseyla7></div> <span class="su-section-code" data-astro-cid-myseyla7>SYS-03</span> <span class="su-section-label" data-astro-cid-myseyla7>TRANSMISSION LOG</span> <div class="su-divider-line" data-astro-cid-myseyla7></div> </div> <h2 class="su-section-title" data-astro-cid-myseyla7>From the field</h2> <p class="su-section-desc" data-astro-cid-myseyla7>Dispatches from startups in the program.</p> <div class="su-tx-grid" data-astro-cid-myseyla7> <div class="su-tx" style="--d:0s" data-astro-cid-myseyla7> <div class="su-tx-header" data-astro-cid-myseyla7> <span class="su-tx-id" data-astro-cid-myseyla7>TX-001</span> <span class="su-tx-from" data-astro-cid-myseyla7>Series A · Healthcare AI</span> </div> <blockquote class="su-tx-msg" data-astro-cid-myseyla7>"We went from 40% hallucination rate to under 2% in our medical Q&amp;A bot. FutureAGI&#39;s evals caught issues our internal tests completely missed."</blockquote> <div class="su-tx-signal" data-astro-cid-myseyla7> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-label" data-astro-cid-myseyla7>SIGNAL STRONG</span> </div> </div><div class="su-tx" style="--d:0.1s" data-astro-cid-myseyla7> <div class="su-tx-header" data-astro-cid-myseyla7> <span class="su-tx-id" data-astro-cid-myseyla7>TX-002</span> <span class="su-tx-from" data-astro-cid-myseyla7>Seed · Legal Tech</span> </div> <blockquote class="su-tx-msg" data-astro-cid-myseyla7>"The startup program gave us enterprise-grade guardrails before we could afford enterprise pricing. Our investors noticed the reliability difference immediately."</blockquote> <div class="su-tx-signal" data-astro-cid-myseyla7> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-label" data-astro-cid-myseyla7>SIGNAL STRONG</span> </div> </div><div class="su-tx" style="--d:0.2s" data-astro-cid-myseyla7> <div class="su-tx-header" data-astro-cid-myseyla7> <span class="su-tx-id" data-astro-cid-myseyla7>TX-003</span> <span class="su-tx-from" data-astro-cid-myseyla7>Pre-seed · EdTech</span> </div> <blockquote class="su-tx-msg" data-astro-cid-myseyla7>"Office hours with the founders helped us redesign our entire eval pipeline. Saved us months of wrong turns."</blockquote> <div class="su-tx-signal" data-astro-cid-myseyla7> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-bar" data-astro-cid-myseyla7></span> <span class="su-signal-label" data-astro-cid-myseyla7>SIGNAL STRONG</span> </div> </div> </div> </div> </section> <!-- ===================== INITIATE LAUNCH: CTA ===================== --> <section id="apply" class="su-section su-launch" data-astro-cid-myseyla7> <div class="su-container" data-astro-cid-myseyla7> <div class="su-launch-box" data-astro-cid-myseyla7> <div class="su-launch-corners" data-astro-cid-myseyla7> <div class="su-lc su-lc-tl" data-astro-cid-myseyla7></div> <div class="su-lc su-lc-tr" data-astro-cid-myseyla7></div> <div class="su-lc su-lc-bl" data-astro-cid-myseyla7></div> <div class="su-lc su-lc-br" data-astro-cid-myseyla7></div> </div> <div class="su-launch-content" data-astro-cid-myseyla7> <div class="su-launch-indicator" data-astro-cid-myseyla7> <span class="su-launch-pulse" data-astro-cid-myseyla7></span> <span class="su-launch-ring" data-astro-cid-myseyla7></span> <span class="su-launch-ring su-launch-ring-2" data-astro-cid-myseyla7></span> </div> <h2 class="su-launch-title" data-astro-cid-myseyla7>Initiate launch sequence</h2> <p class="su-launch-desc" data-astro-cid-myseyla7>
Apply in 2 minutes. Get approved in 48 hours.<br data-astro-cid-myseyla7>
Start building with $6K in credits and full Pro access.
</p> <div class="su-launch-ctas" data-astro-cid-myseyla7> <button id="open-apply-modal" class="su-btn-primary su-btn-lg" data-astro-cid-myseyla7> <span class="su-btn-prompt" data-astro-cid-myseyla7>$</span> launch --now
</button> <a href="https://discord.com/invite/n2tCUKBkAw" target="_blank" rel="noopener noreferrer" class="su-btn-ghost" data-astro-cid-myseyla7> <span class="su-btn-prompt" data-astro-cid-myseyla7>&gt;</span> open_channel --discord
</a> </div> <div class="su-launch-meta" data-astro-cid-myseyla7> <span data-astro-cid-myseyla7>No credit card required</span> <span class="su-meta-dot" data-astro-cid-myseyla7></span> <span data-astro-cid-myseyla7>48hr approval</span> <span class="su-meta-dot" data-astro-cid-myseyla7></span> <span data-astro-cid-myseyla7>Cancel anytime</span> </div> </div> </div> </div> </section> <!-- ===================== FAQ ===================== --> <section class="su-section su-faq" data-astro-cid-myseyla7> <div class="su-container su-container-narrow" data-astro-cid-myseyla7> <div class="su-section-header" data-astro-cid-myseyla7> <div class="su-divider-line" data-astro-cid-myseyla7></div> <span class="su-section-code" data-astro-cid-myseyla7>SYS-04</span> <span class="su-section-label" data-astro-cid-myseyla7>SUPPORT CHANNEL</span> <div class="su-divider-line" data-astro-cid-myseyla7></div> </div> <h2 class="su-section-title" data-astro-cid-myseyla7>Frequently asked</h2> <div class="su-faq-list" data-astro-cid-myseyla7> <div class="su-faq-item" data-faq="0" data-astro-cid-myseyla7> <button class="su-faq-trigger" data-faq-btn="0" data-astro-cid-myseyla7> <span class="su-faq-q" data-astro-cid-myseyla7>What qualifies as a startup?</span> <span class="su-faq-chevron" data-astro-cid-myseyla7> <svg width="14" height="14" viewBox="0 0 14 14" fill="none" stroke="currentColor" stroke-width="1.5" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-myseyla7> <path d="M3.5 5.25L7 8.75L10.5 5.25" data-astro-cid-myseyla7></path> </svg> </span> </button> <div class="su-faq-answer" data-astro-cid-myseyla7> <p data-astro-cid-myseyla7>Companies under $10M raised, less than 5 years old, building with AI. Edge case? Apply anyway. We review every submission individually.</p> </div> </div><div class="su-faq-item" data-faq="1" data-astro-cid-myseyla7> <button class="su-faq-trigger" data-faq-btn="1" data-astro-cid-myseyla7> <span class="su-faq-q" data-astro-cid-myseyla7>How long does approval take?</span> <span class="su-faq-chevron" data-astro-cid-myseyla7> <svg width="14" height="14" viewBox="0 0 14 14" fill="none" stroke="currentColor" stroke-width="1.5" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-myseyla7> <path d="M3.5 5.25L7 8.75L10.5 5.25" data-astro-cid-myseyla7></path> </svg> </span> </button> <div class="su-faq-answer" data-astro-cid-myseyla7> <p data-astro-cid-myseyla7>Most applications are reviewed within 48 hours. You&#39;ll get an email with next steps and your credits will be activated same day.</p> </div> </div><div class="su-faq-item" data-faq="2" data-astro-cid-myseyla7> <button class="su-faq-trigger" data-faq-btn="2" data-astro-cid-myseyla7> <span class="su-faq-q" data-astro-cid-myseyla7>What happens when credits run out?</span> <span class="su-faq-chevron" data-astro-cid-myseyla7> <svg width="14" height="14" viewBox="0 0 14 14" fill="none" stroke="currentColor" stroke-width="1.5" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-myseyla7> <path d="M3.5 5.25L7 8.75L10.5 5.25" data-astro-cid-myseyla7></path> </svg> </span> </button> <div class="su-faq-answer" data-astro-cid-myseyla7> <p data-astro-cid-myseyla7>You transition to standard pricing starting at $99/mo. We also offer extended discounts for program alumni.</p> </div> </div><div class="su-faq-item" data-faq="3" data-astro-cid-myseyla7> <button class="su-faq-trigger" data-faq-btn="3" data-astro-cid-myseyla7> <span class="su-faq-q" data-astro-cid-myseyla7>Which LLM providers are supported?</span> <span class="su-faq-chevron" data-astro-cid-myseyla7> <svg width="14" height="14" viewBox="0 0 14 14" fill="none" stroke="currentColor" stroke-width="1.5" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-myseyla7> <path d="M3.5 5.25L7 8.75L10.5 5.25" data-astro-cid-myseyla7></path> </svg> </span> </button> <div class="su-faq-answer" data-astro-cid-myseyla7> <p data-astro-cid-myseyla7>All of them. OpenAI, Anthropic, Google, Cohere, Mistral, open-source models. Our SDK is provider-agnostic.</p> </div> </div><div class="su-faq-item" data-faq="4" data-astro-cid-myseyla7> <button class="su-faq-trigger" data-faq-btn="4" data-astro-cid-myseyla7> <span class="su-faq-q" data-astro-cid-myseyla7>Can I join if I&#39;m a solo founder?</span> <span class="su-faq-chevron" data-astro-cid-myseyla7> <svg width="14" height="14" viewBox="0 0 14 14" fill="none" stroke="currentColor" stroke-width="1.5" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-myseyla7> <path d="M3.5 5.25L7 8.75L10.5 5.25" data-astro-cid-myseyla7></path> </svg> </span> </button> <div class="su-faq-answer" data-astro-cid-myseyla7> <p data-astro-cid-myseyla7>Absolutely. Solo founders, two-person teams, small crews. The program is designed for early-stage builders of any size.</p> </div> </div> </div> </div> </section> </main>  <div id="apply-modal" class="fixed inset-0 z-[100] opacity-0 pointer-events-none transition-opacity duration-200" data-astro-cid-myseyla7> <div class="absolute inset-0 bg-black/70 backdrop-blur-sm" id="apply-backdrop" data-astro-cid-myseyla7></div> <div class="relative flex items-start justify-center pt-[10vh] px-4" data-astro-cid-myseyla7> <div class="w-full max-w-[520px] bg-[#0a0a0a] border border-[#27272a] rounded-xl shadow-2xl shadow-black/50 overflow-x-hidden flex flex-col max-h-[90vh] transform scale-95 transition-transform duration-200" id="apply-panel" data-astro-cid-myseyla7> <!-- Header --> <div class="flex items-center justify-between px-6 py-4 border-b border-[#1f1f23]" data-astro-cid-myseyla7> <div data-astro-cid-myseyla7> <h3 class="text-lg font-semibold text-white" data-astro-cid-myseyla7>Apply to Startup Program</h3> <p class="text-sm text-zinc-500" data-astro-cid-myseyla7>Takes ~2 minutes. We review within 48 hours.</p> </div> <button id="close-apply-modal" class="w-8 h-8 flex items-center justify-center rounded-lg hover:bg-[#1a1a1a] transition-colors" data-astro-cid-myseyla7> <svg class="w-4 h-4 text-zinc-500" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2" data-astro-cid-myseyla7><path stroke-linecap="round" stroke-linejoin="round" d="M6 18L18 6M6 6l12 12" data-astro-cid-myseyla7></path></svg> </button> </div> <!-- Form --> <form id="startup-apply-form" class="p-6 space-y-4 overflow-y-auto flex-1" data-hs-do-not-collect data-astro-cid-myseyla7> <div class="grid grid-cols-2 gap-4" data-astro-cid-myseyla7> <div data-astro-cid-myseyla7> <label class="block text-sm font-medium text-zinc-400 mb-1.5" data-astro-cid-myseyla7>First name *</label> <input type="text" name="firstName" required class="w-full px-3 py-2.5 bg-[#111111] border border-[#27272a] rounded-lg text-sm text-white placeholder-zinc-600 focus:outline-none focus:border-zinc-500 transition-colors" placeholder="Jane" data-astro-cid-myseyla7> </div> <div data-astro-cid-myseyla7> <label class="block text-sm font-medium text-zinc-400 mb-1.5" data-astro-cid-myseyla7>Last name *</label> <input type="text" name="lastName" required class="w-full px-3 py-2.5 bg-[#111111] border border-[#27272a] rounded-lg text-sm text-white placeholder-zinc-600 focus:outline-none focus:border-zinc-500 transition-colors" placeholder="Doe" data-astro-cid-myseyla7> </div> </div> <div data-astro-cid-myseyla7> <label class="block text-sm font-medium text-zinc-400 mb-1.5" data-astro-cid-myseyla7>Work email *</label> <input type="email" name="email" required class="w-full px-3 py-2.5 bg-[#111111] border border-[#27272a] rounded-lg text-sm text-white placeholder-zinc-600 focus:outline-none focus:border-zinc-500 transition-colors" placeholder="jane@startup.com" data-astro-cid-myseyla7> </div> <div data-astro-cid-myseyla7> <label class="block text-sm font-medium text-zinc-400 mb-1.5" data-astro-cid-myseyla7>Company name *</label> <input type="text" name="company" required class="w-full px-3 py-2.5 bg-[#111111] border border-[#27272a] rounded-lg text-sm text-white placeholder-zinc-600 focus:outline-none focus:border-zinc-500 transition-colors" placeholder="Acme AI" data-astro-cid-myseyla7> </div> <div data-astro-cid-myseyla7> <label class="block text-sm font-medium text-zinc-400 mb-1.5" data-astro-cid-myseyla7>Company website</label> <input type="url" name="website" class="w-full px-3 py-2.5 bg-[#111111] border border-[#27272a] rounded-lg text-sm text-white placeholder-zinc-600 focus:outline-none focus:border-zinc-500 transition-colors" placeholder="https://acme.ai" data-astro-cid-myseyla7> </div> <div data-astro-cid-myseyla7> <label class="block text-sm font-medium text-zinc-400 mb-1.5" data-astro-cid-myseyla7>Stage *</label> <select name="stage" required class="w-full px-3 py-2.5 bg-[#111111] border border-[#27272a] rounded-lg text-sm text-white focus:outline-none focus:border-zinc-500 transition-colors appearance-none" data-astro-cid-myseyla7> <option value="" class="text-zinc-600" data-astro-cid-myseyla7>Select your stage</option> <option value="pre-seed" data-astro-cid-myseyla7>Pre-seed</option> <option value="seed" data-astro-cid-myseyla7>Seed</option> <option value="series-a" data-astro-cid-myseyla7>Series A</option> <option value="series-b" data-astro-cid-myseyla7>Series B</option> <option value="bootstrapped" data-astro-cid-myseyla7>Bootstrapped</option> </select> </div> <div data-astro-cid-myseyla7> <label class="block text-sm font-medium text-zinc-400 mb-1.5" data-astro-cid-myseyla7>What are you building? *</label> <textarea name="description" required rows="3" class="w-full px-3 py-2.5 bg-[#111111] border border-[#27272a] rounded-lg text-sm text-white placeholder-zinc-600 focus:outline-none focus:border-zinc-500 transition-colors resize-none" placeholder="Tell us about your AI product and how you plan to use Future AGI..." data-astro-cid-myseyla7></textarea> </div> <!-- Submit --> <button type="submit" id="apply-submit" class="w-full py-3 bg-white text-black text-sm font-medium rounded-lg hover:bg-zinc-200 transition-colors" data-astro-cid-myseyla7>
Submit Application
</button> <p class="text-xs text-zinc-600 text-center" data-astro-cid-myseyla7>
No credit card required. We review every application personally.
</p> </form> <!-- Success state outside form so it isn't hidden by form querySelectorAll --> <div id="apply-success" class="hidden text-center py-8 px-6" data-astro-cid-myseyla7> <div class="w-14 h-14 mx-auto mb-4 rounded-full bg-[#22c55e]/10 flex items-center justify-center" data-astro-cid-myseyla7> <svg class="w-7 h-7 text-[#22c55e]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2" data-astro-cid-myseyla7><path stroke-linecap="round" stroke-linejoin="round" d="M4.5 12.75l6 6 9-13.5" data-astro-cid-myseyla7></path></svg> </div> <h4 class="text-lg font-semibold text-white mb-2" data-astro-cid-myseyla7>Application received!</h4> <p class="text-sm text-zinc-400" data-astro-cid-myseyla7>We'll review your application and get back to you within 48 hours.</p> </div> </div> </div> </div> <footer class="relative bg-[#0a0a0a] border-t border-[#27272a] overflow-hidden" data-astro-cid-35ed7um5> <!-- Launch Ready Section - Full Width Starship Banner --> <div class="launch-ready-section relative py-16 lg:py-24 border-b border-[#27272a]" data-astro-cid-35ed7um5> <!-- Starfield Background - Jump sequence then new galaxy reveal --> <div class="absolute inset-0 overflow-hidden pointer-events-none" data-astro-cid-35ed7um5> <!-- Original stars - streak during jump, then disappear --> <div class="star star-1" data-astro-cid-35ed7um5></div> <div class="star star-2" data-astro-cid-35ed7um5></div> <div class="star star-3" data-astro-cid-35ed7um5></div> <div class="star star-4" data-astro-cid-35ed7um5></div> <div class="star star-5" data-astro-cid-35ed7um5></div> <div class="star star-6" data-astro-cid-35ed7um5></div> <div class="star star-7" data-astro-cid-35ed7um5></div> <div class="star star-8" data-astro-cid-35ed7um5></div> <div class="star star-9" data-astro-cid-35ed7um5></div> <div class="star star-10" data-astro-cid-35ed7um5></div> <div class="star star-11" data-astro-cid-35ed7um5></div> <div class="star star-12" data-astro-cid-35ed7um5></div> <!-- Hyperspace streaks - play once during jump --> <div class="hyperspace-streak streak-1" data-astro-cid-35ed7um5></div> <div class="hyperspace-streak streak-2" data-astro-cid-35ed7um5></div> <div class="hyperspace-streak streak-3" data-astro-cid-35ed7um5></div> <div class="hyperspace-streak streak-4" data-astro-cid-35ed7um5></div> <div class="hyperspace-streak streak-5" data-astro-cid-35ed7um5></div> <!-- NEW GALAXY - Fades in after jump completes --> <div class="new-galaxy" data-astro-cid-35ed7um5> <!-- Distant nebula glow --> <div class="nebula nebula-1" data-astro-cid-35ed7um5></div> <div class="nebula nebula-2" data-astro-cid-35ed7um5></div> <div class="nebula nebula-3" data-astro-cid-35ed7um5></div> <!-- Slowly drifting celestial bodies - parallax layers --> <!-- Far background planets (slow) --> <div class="celestial-body planet-far-1" data-astro-cid-35ed7um5></div> <div class="celestial-body planet-far-2" data-astro-cid-35ed7um5></div> <div class="celestial-body planet-far-3" data-astro-cid-35ed7um5></div> <!-- Mid-ground planets (medium speed) --> <div class="celestial-body planet-mid-1" data-astro-cid-35ed7um5></div> <div class="celestial-body planet-mid-2" data-astro-cid-35ed7um5></div> <div class="celestial-body planet-mid-3" data-astro-cid-35ed7um5></div> <!-- Closer moons (faster) --> <div class="celestial-body moon-1" data-astro-cid-35ed7um5></div> <div class="celestial-body moon-2" data-astro-cid-35ed7um5></div> <div class="celestial-body moon-3" data-astro-cid-35ed7um5></div> <div class="celestial-body moon-4" data-astro-cid-35ed7um5></div> <!-- Asteroids (fastest - closest) --> <div class="celestial-body asteroid-1" data-astro-cid-35ed7um5></div> <div class="celestial-body asteroid-2" data-astro-cid-35ed7um5></div> <div class="celestial-body asteroid-3" data-astro-cid-35ed7um5></div> <div class="celestial-body asteroid-4" data-astro-cid-35ed7um5></div> <div class="celestial-body asteroid-5" data-astro-cid-35ed7um5></div> <!-- Background stars - far layer (slow drift) --> <div class="galaxy-star star-far sf-1" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-2" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-3" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-4" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-5" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-6" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-7" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-8" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-9" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-10" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-11" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-far sf-12" data-astro-cid-35ed7um5></div> <!-- Mid-ground stars (medium drift) --> <div class="galaxy-star star-mid sm-1" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-2" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-3" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-4" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-5" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-6" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-7" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-8" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-9" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-mid sm-10" data-astro-cid-35ed7um5></div> <!-- Close stars (faster drift) --> <div class="galaxy-star star-close sc-1" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-close sc-2" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-close sc-3" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-close sc-4" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-close sc-5" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-close sc-6" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-close sc-7" data-astro-cid-35ed7um5></div> <div class="galaxy-star star-close sc-8" data-astro-cid-35ed7um5></div> </div> </div> <!-- Subtle grid overlay --> <div class="absolute inset-0 opacity-[0.03]" style="background-image: linear-gradient(#fafafa 1px, transparent 1px), linear-gradient(90deg, #fafafa 1px, transparent 1px); background-size: 40px 40px;" data-astro-cid-35ed7um5></div> <div class="max-w-[1400px] mx-auto px-6 lg:px-8 relative z-10" data-astro-cid-35ed7um5> <div class="flex flex-col lg:flex-row items-center justify-between gap-12 lg:gap-16" data-astro-cid-35ed7um5> <!-- Left Side - Starship with Hyperspace Effect --> <div class="starship-container flex-1 flex justify-center lg:justify-start overflow-visible" data-astro-cid-35ed7um5> <svg class="starship-svg w-[300px] h-[160px] md:w-[380px] md:h-[200px] overflow-visible" viewBox="0 0 420 200" fill="none" xmlns="http://www.w3.org/2000/svg" style="overflow: visible;" data-astro-cid-35ed7um5> <!-- Clean horizontal starship - pointing right --> <g class="ship" data-astro-cid-35ed7um5> <!-- Engine exhaust glow --> <g class="engine-glow" data-astro-cid-35ed7um5> <ellipse cx="45" cy="100" rx="20" ry="6" fill="url(#exhaustGlow)" opacity="0.8" data-astro-cid-35ed7um5></ellipse> <ellipse cx="35" cy="100" rx="30" ry="10" fill="url(#exhaustGlow)" opacity="0.4" data-astro-cid-35ed7um5></ellipse> <ellipse cx="25" cy="100" rx="40" ry="14" fill="url(#exhaustGlow)" opacity="0.2" data-astro-cid-35ed7um5></ellipse> </g> <!-- Main fuselage --> <path d="M60 85 L60 115 L270 118 L270 82 L60 85 Z" fill="#18181b" stroke="#52525b" stroke-width="1.5" class="fuselage" data-astro-cid-35ed7um5></path> <!-- Nose cone - sleek pointed design --> <path d="M270 82 L270 118 L310 115 L355 100 L310 85 L270 82 Z" fill="#1f1f23" stroke="#52525b" stroke-width="1" data-astro-cid-35ed7um5></path> <!-- Nose cone highlight (top surface catching light) --> <path d="M270 82 L310 85 L355 100 L310 92 L270 88 Z" fill="#27272a" opacity="0.6" data-astro-cid-35ed7um5></path> <!-- Nose cone shadow (bottom surface) --> <path d="M270 118 L310 115 L355 100 L310 108 L270 112 Z" fill="#141417" opacity="0.8" data-astro-cid-35ed7um5></path> <!-- Nose tip accent --> <path d="M340 94 L355 100 L340 106" fill="none" stroke="#71717a" stroke-width="0.5" data-astro-cid-35ed7um5></path> <!-- Cockpit canopy frame --> <path d="M275 88 L275 112 Q285 115 295 112 L320 105 L320 95 L295 88 Q285 85 275 88 Z" fill="#0f0f12" stroke="#52525b" stroke-width="1" data-astro-cid-35ed7um5></path> <!-- Cockpit glass - outer --> <path d="M280 90 L280 110 Q288 112 296 109 L315 103 L315 97 L296 91 Q288 88 280 90 Z" fill="url(#cockpitGlass)" stroke="#3f3f46" stroke-width="0.5" data-astro-cid-35ed7um5></path> <!-- Cockpit glass reflection --> <path d="M282 91 L282 100 Q288 101 294 99 L308 96 L308 94 L294 92 Q288 90 282 91 Z" fill="#fafafa" opacity="0.08" data-astro-cid-35ed7um5></path> <!-- Cockpit internal frame lines --> <line x1="290" y1="90" x2="290" y2="110" stroke="#27272a" stroke-width="0.5" data-astro-cid-35ed7um5></line> <line x1="302" y1="92" x2="302" y2="108" stroke="#27272a" stroke-width="0.5" data-astro-cid-35ed7um5></line> <!-- Panel lines on nose --> <line x1="270" y1="95" x2="330" y2="97" stroke="#3f3f46" stroke-width="0.3" data-astro-cid-35ed7um5></line> <line x1="270" y1="105" x2="330" y2="103" stroke="#3f3f46" stroke-width="0.3" data-astro-cid-35ed7um5></line> <!-- Upper wing --> <path d="M100 85 L80 45 L140 45 L160 80" fill="#1a1a1a" stroke="#52525b" stroke-width="1" data-astro-cid-35ed7um5></path> <line x1="95" y1="65" x2="135" y2="65" stroke="#3f3f46" stroke-width="0.5" data-astro-cid-35ed7um5></line> <!-- Lower wing --> <path d="M100 115 L80 155 L140 155 L160 120" fill="#1a1a1a" stroke="#52525b" stroke-width="1" data-astro-cid-35ed7um5></path> <line x1="95" y1="135" x2="135" y2="135" stroke="#3f3f46" stroke-width="0.5" data-astro-cid-35ed7um5></line> <!-- Engine block --> <rect x="55" y="88" width="30" height="24" rx="2" fill="#1a1a1a" stroke="#52525b" data-astro-cid-35ed7um5></rect> <rect x="58" y="92" width="8" height="16" rx="1" fill="#27272a" stroke="#3f3f46" data-astro-cid-35ed7um5></rect> <!-- Body details --> <line x1="120" y1="85" x2="120" y2="115" stroke="#3f3f46" stroke-width="0.5" data-astro-cid-35ed7um5></line> <line x1="180" y1="83" x2="180" y2="117" stroke="#3f3f46" stroke-width="0.5" data-astro-cid-35ed7um5></line> <line x1="240" y1="82" x2="240" y2="118" stroke="#3f3f46" stroke-width="0.5" data-astro-cid-35ed7um5></line> <!-- Antenna --> <line x1="355" y1="100" x2="370" y2="100" stroke="#71717a" stroke-width="1" data-astro-cid-35ed7um5></line> <circle cx="373" cy="100" r="2.5" fill="#27272a" stroke="#71717a" stroke-width="0.5" data-astro-cid-35ed7um5></circle> <!-- Navigation lights --> <circle cx="78" cy="45" r="2" fill="#a1a1aa" class="nav-light" data-astro-cid-35ed7um5> <animate attributeName="opacity" values="1;0.3;1" dur="2s" repeatCount="indefinite" data-astro-cid-35ed7um5></animate> </circle> <circle cx="78" cy="155" r="2" fill="#a1a1aa" class="nav-light" data-astro-cid-35ed7um5> <animate attributeName="opacity" values="1;0.3;1" dur="2s" repeatCount="indefinite" begin="0.5s" data-astro-cid-35ed7um5></animate> </circle> <!-- Front nav light on nose tip --> <circle cx="355" cy="100" r="1.5" fill="#fafafa" class="nav-light" data-astro-cid-35ed7um5> <animate attributeName="opacity" values="1;0.4;1" dur="1.2s" repeatCount="indefinite" data-astro-cid-35ed7um5></animate> </circle> </g> <!-- Gradients --> <defs data-astro-cid-35ed7um5> <radialGradient id="exhaustGlow" cx="100%" cy="50%" r="100%" data-astro-cid-35ed7um5> <stop offset="0%" stop-color="#fafafa" data-astro-cid-35ed7um5></stop> <stop offset="50%" stop-color="#a1a1aa" data-astro-cid-35ed7um5></stop> <stop offset="100%" stop-color="transparent" data-astro-cid-35ed7um5></stop> </radialGradient> <!-- Cockpit glass gradient - dark with subtle blue-ish tint --> <linearGradient id="cockpitGlass" x1="0%" y1="0%" x2="100%" y2="100%" data-astro-cid-35ed7um5> <stop offset="0%" stop-color="#1a1a1f" data-astro-cid-35ed7um5></stop> <stop offset="50%" stop-color="#0f0f14" data-astro-cid-35ed7um5></stop> <stop offset="100%" stop-color="#18181d" data-astro-cid-35ed7um5></stop> </linearGradient> </defs> </svg> </div> <!-- Right Side - Status & CTA --> <div class="flex-1 text-center lg:text-left max-w-md" data-astro-cid-35ed7um5> <!-- Status Badge --> <div class="inline-flex items-center gap-2 px-3 py-1.5 rounded-full bg-[#fafafa]/5 border border-[#3f3f46] mb-4" data-astro-cid-35ed7um5> <span class="relative flex h-2 w-2" data-astro-cid-35ed7um5> <span class="animate-ping absolute inline-flex h-full w-full rounded-full bg-[#fafafa] opacity-50" data-astro-cid-35ed7um5></span> <span class="relative inline-flex rounded-full h-2 w-2 bg-[#fafafa]" data-astro-cid-35ed7um5></span> </span> <span class="text-[12px] text-[#fafafa] font-medium uppercase tracking-wider" data-astro-cid-35ed7um5>Ready for Launch</span> </div> <!-- Headline --> <div class="text-2xl md:text-3xl lg:text-4xl font-semibold text-[#fafafa] mb-3 tracking-tight" data-astro-cid-35ed7um5>
Your agent is<br class="hidden sm:block" data-astro-cid-35ed7um5> ready for launch
</div> <!-- Phase Checklist --> <div class="flex flex-wrap justify-center lg:justify-start gap-2 mb-6" data-astro-cid-35ed7um5> <div class="phase-complete inline-flex items-center gap-1.5 px-2.5 py-1 rounded bg-[#18181b] border border-[#3f3f46]" data-astro-cid-35ed7um5> <svg class="w-3.5 h-3.5 text-[#fafafa]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2.5" data-astro-cid-35ed7um5> <path stroke-linecap="round" stroke-linejoin="round" d="M5 13l4 4L19 7" data-astro-cid-35ed7um5></path> </svg> <span class="text-[11px] text-[#a1a1aa] uppercase tracking-wider" data-astro-cid-35ed7um5>Prompts</span> </div> <div class="phase-complete inline-flex items-center gap-1.5 px-2.5 py-1 rounded bg-[#18181b] border border-[#3f3f46]" data-astro-cid-35ed7um5> <svg class="w-3.5 h-3.5 text-[#fafafa]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2.5" data-astro-cid-35ed7um5> <path stroke-linecap="round" stroke-linejoin="round" d="M5 13l4 4L19 7" data-astro-cid-35ed7um5></path> </svg> <span class="text-[11px] text-[#a1a1aa] uppercase tracking-wider" data-astro-cid-35ed7um5>Simulate</span> </div> <div class="phase-complete inline-flex items-center gap-1.5 px-2.5 py-1 rounded bg-[#18181b] border border-[#3f3f46]" data-astro-cid-35ed7um5> <svg class="w-3.5 h-3.5 text-[#fafafa]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2.5" data-astro-cid-35ed7um5> <path stroke-linecap="round" stroke-linejoin="round" d="M5 13l4 4L19 7" data-astro-cid-35ed7um5></path> </svg> <span class="text-[11px] text-[#a1a1aa] uppercase tracking-wider" data-astro-cid-35ed7um5>Evaluate</span> </div> <div class="phase-complete inline-flex items-center gap-1.5 px-2.5 py-1 rounded bg-[#18181b] border border-[#3f3f46]" data-astro-cid-35ed7um5> <svg class="w-3.5 h-3.5 text-[#fafafa]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2.5" data-astro-cid-35ed7um5> <path stroke-linecap="round" stroke-linejoin="round" d="M5 13l4 4L19 7" data-astro-cid-35ed7um5></path> </svg> <span class="text-[11px] text-[#a1a1aa] uppercase tracking-wider" data-astro-cid-35ed7um5>Optimize</span> </div> <div class="phase-complete inline-flex items-center gap-1.5 px-2.5 py-1 rounded bg-[#18181b] border border-[#3f3f46]" data-astro-cid-35ed7um5> <svg class="w-3.5 h-3.5 text-[#fafafa]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2.5" data-astro-cid-35ed7um5> <path stroke-linecap="round" stroke-linejoin="round" d="M5 13l4 4L19 7" data-astro-cid-35ed7um5></path> </svg> <span class="text-[11px] text-[#a1a1aa] uppercase tracking-wider" data-astro-cid-35ed7um5>Observe</span> </div> <div class="phase-complete inline-flex items-center gap-1.5 px-2.5 py-1 rounded bg-[#18181b] border border-[#3f3f46]" data-astro-cid-35ed7um5> <svg class="w-3.5 h-3.5 text-[#fafafa]" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2.5" data-astro-cid-35ed7um5> <path stroke-linecap="round" stroke-linejoin="round" d="M5 13l4 4L19 7" data-astro-cid-35ed7um5></path> </svg> <span class="text-[11px] text-[#a1a1aa] uppercase tracking-wider" data-astro-cid-35ed7um5>Guard</span> </div> </div> <!-- Tagline --> <p class="text-[15px] text-[#d4d4d8] mb-6" data-astro-cid-35ed7um5>
Test, guard, and monitor complete. Ship with confidence.
</p> <!-- CTA Buttons --> <div class="flex flex-col sm:flex-row gap-3 justify-center lg:justify-start" data-astro-cid-35ed7um5> <a href="https://app.futureagi.com/auth/jwt/register" target="_blank" rel="noopener noreferrer" class="launch-button group inline-flex items-center justify-center gap-2 px-6 py-3 bg-[#fafafa] text-[#0a0a0a] text-[14px] font-medium rounded-full hover:bg-white transition-all hover:shadow-[0_0_30px_rgba(255,255,255,0.3)]" data-astro-cid-35ed7um5> <svg class="w-4 h-4 transition-transform group-hover:-translate-y-0.5 group-hover:translate-x-0.5" fill="none" viewBox="0 0 24 24" stroke="currentColor" stroke-width="2" data-astro-cid-35ed7um5> <path stroke-linecap="round" stroke-linejoin="round" d="M4.5 19.5l15-15m0 0H8.25m11.25 0v11.25" data-astro-cid-35ed7um5></path> </svg>
Deploy to Production
</a> <a href="https://docs.futureagi.com" target="_blank" rel="noopener noreferrer" class="inline-flex items-center justify-center gap-2 px-6 py-3 text-[#a1a1aa] text-[14px] font-medium rounded-full border border-[#27272a] hover:border-[#3f3f46] hover:text-[#fafafa] transition-all" data-astro-cid-35ed7um5>
View Documentation
</a> </div> </div> </div> </div> </div> <div class="max-w-[1400px] mx-auto px-6 lg:px-8" data-astro-cid-35ed7um5> <!-- Main Footer Content --> <div class="py-12 lg:py-16" data-astro-cid-35ed7um5> <!-- Link Columns --> <div class="grid grid-cols-2 md:grid-cols-3 lg:grid-cols-6 gap-8 lg:gap-6" data-astro-cid-35ed7um5> <!-- Platform --> <div data-astro-cid-35ed7um5> <div class="text-[12px] text-[#fafafa] font-semibold tracking-wide uppercase mb-4" data-astro-cid-35ed7um5>Platform</div> <ul class="space-y-3" data-astro-cid-35ed7um5> <li data-astro-cid-35ed7um5> <a href="/platform/simulate/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Simulate  </a> </li><li data-astro-cid-35ed7um5> <a href="/platform/evaluate/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Evaluate  </a> </li><li data-astro-cid-35ed7um5> <a href="/platform/guard/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Guard <span class="px-1.5 py-0.5 text-[9px] font-semibold bg-[#fafafa] text-[#0a0a0a] rounded uppercase tracking-wide" data-astro-cid-35ed7um5>New</span> </a> </li><li data-astro-cid-35ed7um5> <a href="/platform/monitor/tracing/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Monitor  </a> </li><li data-astro-cid-35ed7um5> <a href="/platform/optimize/rl/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Optimize  </a> </li><li data-astro-cid-35ed7um5> <a href="/platform/agents/ide/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Agents  </a> </li> </ul> </div> <!-- Solutions --> <div data-astro-cid-35ed7um5> <div class="text-[12px] text-[#fafafa] font-semibold tracking-wide uppercase mb-4" data-astro-cid-35ed7um5>Solutions</div> <ul class="space-y-3" data-astro-cid-35ed7um5> <li data-astro-cid-35ed7um5> <a href="/enterprise/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Enterprise  </a> </li><li data-astro-cid-35ed7um5> <a href="/startups/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Startups  </a> </li><li data-astro-cid-35ed7um5> <a href="/non-profit/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Non-Profit  </a> </li><li data-astro-cid-35ed7um5> <a href="/pricing/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Pricing  </a> </li><li data-astro-cid-35ed7um5> <a href="/watch-demo/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Watch Demo  </a> </li><li data-astro-cid-35ed7um5> <a href="/talk-to-human/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors inline-flex items-center gap-2" data-astro-cid-35ed7um5> Talk to a Human  </a> </li> </ul> </div> <!-- Resources --> <div data-astro-cid-35ed7um5> <div class="text-[12px] text-[#fafafa] font-semibold tracking-wide uppercase mb-4" data-astro-cid-35ed7um5>Resources</div> <ul class="space-y-3" data-astro-cid-35ed7um5> <li data-astro-cid-35ed7um5> <a href="/blog/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Blog </a> </li><li data-astro-cid-35ed7um5> <a href="/glossary/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Glossary </a> </li><li data-astro-cid-35ed7um5> <a href="/llm-cost-calculator/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> LLM Cost Calculator </a> </li><li data-astro-cid-35ed7um5> <a href="/eval-tco-calculator/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Evaluation TCO Calculator </a> </li><li data-astro-cid-35ed7um5> <a href="/customers/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Customer Stories </a> </li><li data-astro-cid-35ed7um5> <a href="/research/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Research </a> </li><li data-astro-cid-35ed7um5> <a href="/ebooks/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> eBooks </a> </li><li data-astro-cid-35ed7um5> <a href="/handbook/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Handbook </a> </li><li data-astro-cid-35ed7um5> <a href="/changelog/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Changelog </a> </li><li data-astro-cid-35ed7um5> <a href="/roadmap/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Roadmap </a> </li> </ul> </div> <!-- Developers --> <div data-astro-cid-35ed7um5> <div class="text-[12px] text-[#fafafa] font-semibold tracking-wide uppercase mb-4" data-astro-cid-35ed7um5>Developers</div> <ul class="space-y-3" data-astro-cid-35ed7um5> <li data-astro-cid-35ed7um5> <a href="https://docs.futureagi.com" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> Documentation </a> </li><li data-astro-cid-35ed7um5> <a href="https://docs.futureagi.com/docs/quickstart/setup-observability" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> Quick Start </a> </li><li data-astro-cid-35ed7um5> <a href="https://docs.futureagi.com/docs/api" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> API Reference </a> </li><li data-astro-cid-35ed7um5> <a href="https://docs.futureagi.com/docs/sdk" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> SDKs </a> </li><li data-astro-cid-35ed7um5> <a href="/integrations/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Integrations </a> </li><li data-astro-cid-35ed7um5> <a href="https://github.com/future-agi" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> GitHub </a> </li><li data-astro-cid-35ed7um5> <a href="https://status.futureagi.com" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> Status </a> </li> </ul> </div> <!-- Company --> <div data-astro-cid-35ed7um5> <div class="text-[12px] text-[#fafafa] font-semibold tracking-wide uppercase mb-4" data-astro-cid-35ed7um5>Company</div> <ul class="space-y-3" data-astro-cid-35ed7um5> <li data-astro-cid-35ed7um5> <a href="/careers/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Careers </a> </li><li data-astro-cid-35ed7um5> <a href="/contact/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Contact </a> </li><li data-astro-cid-35ed7um5> <a href="/security/" class="text-[14px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Security </a> </li> </ul> </div> <!-- Connect (Social + Newsletter) --> <div data-astro-cid-35ed7um5> <div class="text-[12px] text-[#fafafa] font-semibold tracking-wide uppercase mb-4" data-astro-cid-35ed7um5>Connect</div> <div class="flex items-center gap-1 mb-6" data-astro-cid-35ed7um5> <a href="https://www.linkedin.com/company/futureagi/" class="w-9 h-9 flex items-center justify-center text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#27272a] rounded-lg transition-all" aria-label="LinkedIn" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> <svg class="w-[18px] h-[18px]" fill="currentColor" viewBox="0 0 24 24" data-astro-cid-35ed7um5> <path d="M20.447 20.452h-3.554v-5.569c0-1.328-.027-3.037-1.852-3.037-1.853 0-2.136 1.445-2.136 2.939v5.667H9.351V9h3.414v1.561h.046c.477-.9 1.637-1.85 3.37-1.85 3.601 0 4.267 2.37 4.267 5.455v6.286zM5.337 7.433c-1.144 0-2.063-.926-2.063-2.065 0-1.138.92-2.063 2.063-2.063 1.14 0 2.064.925 2.064 2.063 0 1.139-.925 2.065-2.064 2.065zm1.782 13.019H3.555V9h3.564v11.452zM22.225 0H1.771C.792 0 0 .774 0 1.729v20.542C0 23.227.792 24 1.771 24h20.451C23.2 24 24 23.227 24 22.271V1.729C24 .774 23.2 0 22.222 0h.003z"/> </svg> </a><a href="https://x.com/FutureAGI_" class="w-9 h-9 flex items-center justify-center text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#27272a] rounded-lg transition-all" aria-label="X" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> <svg class="w-[18px] h-[18px]" fill="currentColor" viewBox="0 0 24 24" data-astro-cid-35ed7um5> <path d="M18.244 2.25h3.308l-7.227 8.26 8.502 11.24H16.17l-5.214-6.817L4.99 21.75H1.68l7.73-8.835L1.254 2.25H8.08l4.713 6.231zm-1.161 17.52h1.833L7.084 4.126H5.117z"/> </svg> </a><a href="https://discord.com/invite/n2tCUKBkAw" class="w-9 h-9 flex items-center justify-center text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#27272a] rounded-lg transition-all" aria-label="Discord" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> <svg class="w-[18px] h-[18px]" fill="currentColor" viewBox="0 0 24 24" data-astro-cid-35ed7um5> <path d="M20.317 4.37a19.791 19.791 0 0 0-4.885-1.515.074.074 0 0 0-.079.037c-.21.375-.444.864-.608 1.25a18.27 18.27 0 0 0-5.487 0 12.64 12.64 0 0 0-.617-1.25.077.077 0 0 0-.079-.037A19.736 19.736 0 0 0 3.677 4.37a.07.07 0 0 0-.032.027C.533 9.046-.32 13.58.099 18.057a.082.082 0 0 0 .031.057 19.9 19.9 0 0 0 5.993 3.03.078.078 0 0 0 .084-.028 14.09 14.09 0 0 0 1.226-1.994.076.076 0 0 0-.041-.106 13.107 13.107 0 0 1-1.872-.892.077.077 0 0 1-.008-.128 10.2 10.2 0 0 0 .372-.292.074.074 0 0 1 .077-.01c3.928 1.793 8.18 1.793 12.062 0a.074.074 0 0 1 .078.01c.12.098.246.198.373.292a.077.077 0 0 1-.006.127 12.299 12.299 0 0 1-1.873.892.077.077 0 0 0-.041.107c.36.698.772 1.362 1.225 1.993a.076.076 0 0 0 .084.028 19.839 19.839 0 0 0 6.002-3.03.077.077 0 0 0 .032-.054c.5-5.177-.838-9.674-3.549-13.66a.061.061 0 0 0-.031-.03zM8.02 15.33c-1.183 0-2.157-1.085-2.157-2.419 0-1.333.956-2.419 2.157-2.419 1.21 0 2.176 1.096 2.157 2.42 0 1.333-.956 2.418-2.157 2.418zm7.975 0c-1.183 0-2.157-1.085-2.157-2.419 0-1.333.955-2.419 2.157-2.419 1.21 0 2.176 1.096 2.157 2.42 0 1.333-.946 2.418-2.157 2.418z"/> </svg> </a><a href="https://github.com/future-agi" class="w-9 h-9 flex items-center justify-center text-[#a1a1aa] hover:text-[#fafafa] hover:bg-[#27272a] rounded-lg transition-all" aria-label="GitHub" target="_blank" rel="noopener noreferrer" data-astro-cid-35ed7um5> <svg class="w-[18px] h-[18px]" fill="currentColor" viewBox="0 0 24 24" data-astro-cid-35ed7um5> <path d="M12 0c-6.626 0-12 5.373-12 12 0 5.302 3.438 9.8 8.207 11.387.599.111.793-.261.793-.577v-2.234c-3.338.726-4.033-1.416-4.033-1.416-.546-1.387-1.333-1.756-1.333-1.756-1.089-.745.083-.729.083-.729 1.205.084 1.839 1.237 1.839 1.237 1.07 1.834 2.807 1.304 3.492.997.107-.775.418-1.305.762-1.604-2.665-.305-5.467-1.334-5.467-5.931 0-1.311.469-2.381 1.236-3.221-.124-.303-.535-1.524.117-3.176 0 0 1.008-.322 3.301 1.23.957-.266 1.983-.399 3.003-.404 1.02.005 2.047.138 3.006.404 2.291-1.552 3.297-1.23 3.297-1.23.653 1.653.242 2.874.118 3.176.77.84 1.235 1.911 1.235 3.221 0 4.609-2.807 5.624-5.479 5.921.43.372.823 1.102.823 2.222v3.293c0 .319.192.694.801.576 4.765-1.589 8.199-6.086 8.199-11.386 0-6.627-5.373-12-12-12z"/> </svg> </a> </div> </div> </div> </div> <!-- Bottom Section --> <div class="py-6 border-t border-[#27272a]" data-astro-cid-35ed7um5> <div class="flex flex-col lg:flex-row items-start lg:items-center justify-between gap-6" data-astro-cid-35ed7um5> <!-- Logo + Copyright --> <div class="flex flex-col sm:flex-row items-start sm:items-center gap-4 sm:gap-6" data-astro-cid-35ed7um5> <div class="inline-flex items-center gap-1.5"><svg class="h-[18px] w-auto flex-shrink-0" viewBox="0 0 47 47" fill="none" xmlns="http://www.w3.org/2000/svg"><path d="M46.8996 25.4157L43.4957 27.3812L40.0896 29.3467L36.6856 31.3143L33.2816 33.2798L31.314 36.686L29.3485 40.0899L27.383 43.496L25.4175 46.9H21.4844L23.4499 43.496L25.4175 40.0899L27.383 36.686L29.3485 33.2798L30.787 30.7852L33.2795 29.3467L36.6856 27.3812L40.0896 25.4157L36.6856 23.4502L33.2795 21.4826L29.8733 19.5171L28.2923 18.6056L27.3808 17.0268L25.4153 13.6206L23.4499 10.2145L25.4153 6.81055L27.3808 10.2145L29.3463 13.6206L30.7848 16.1131L33.2795 17.5516L36.6856 19.5171L40.0896 21.4826L43.4957 23.4502L46.8996 25.4157Z" fill="url(#star-g1)"></path><path d="M40.0895 25.4153L36.6855 27.3808L33.2794 29.3463L30.7869 30.7848L29.3484 33.2795L27.3829 36.6856L25.4174 40.0896L23.4498 43.4957L21.4843 46.8996L19.5188 43.4957L17.5533 40.0896L15.5857 36.6856L13.6202 33.2816L10.2141 31.314L6.81009 29.3485L3.40397 27.383L0 25.4175V21.4844L3.40397 23.4499L6.81009 25.4175L10.2141 27.383L13.6202 29.3485L16.1148 30.787L17.5533 33.2795L19.5188 36.6856L21.4843 40.0896L23.4498 36.6856L25.4174 33.2795L27.3829 29.8733L28.2944 28.2923L29.8732 27.3808L33.2794 25.4153L36.6855 23.4499L40.0895 25.4153Z" fill="url(#star-g2)"></path><path d="M10.2141 19.5188L6.81009 21.4843L10.2141 23.4498L13.6202 25.4174L17.0263 27.3829L18.6073 28.2944L19.5188 29.8732L21.4843 33.2794L23.4498 36.6855L21.4843 40.0895L19.5188 36.6855L17.5533 33.2794L16.1148 30.7869L13.6202 29.3484L10.2141 27.3829L6.81009 25.4174L3.40397 23.4498L0 21.4843L3.40397 19.5188L6.81009 17.5533L10.2141 15.5857L13.618 13.6202L15.5857 10.2141L17.5512 6.81009L19.5166 3.40397L21.4821 0H25.4153L23.4498 3.40397L21.4821 6.81009L19.5166 10.2141L17.5512 13.6202L16.1127 16.1148L15.2638 16.6051L13.6202 17.5533L10.2141 19.5188Z" fill="url(#star-g3)"></path><path d="M46.9 21.4821V25.4153L43.496 23.4498L40.0899 21.4821L36.686 19.5166L33.2798 17.5512L30.7852 16.1127L29.3467 13.6202L27.3812 10.2141L25.4157 6.81009L23.4502 10.2141L21.4826 13.6202L19.5171 17.0263L18.6056 18.6073L17.0268 19.5188L13.6206 21.4843L10.2145 23.4498L6.81055 21.4843L10.2145 19.5188L13.6206 17.5533L15.2643 16.6051L16.1131 16.1148L17.5516 13.6202L19.5171 10.2141L21.4826 6.81009L23.4502 3.40397L25.4157 0L27.3812 3.40397L29.3467 6.81009L31.3143 10.2141L33.2798 13.618L36.686 15.5857L40.0899 17.5512L43.496 19.5166L46.9 21.4821Z" fill="url(#star-g4)"></path><defs><linearGradient id="star-g1" x1="34.192" y1="6.81055" x2="34.192" y2="46.9" gradientUnits="userSpaceOnUse"><stop stop-color="white"></stop><stop offset="1" stop-color="#E6E6E7"></stop></linearGradient><linearGradient id="star-g2" x1="20.0447" y1="21.4844" x2="20.0447" y2="46.8996" gradientUnits="userSpaceOnUse"><stop stop-color="#F3F3F3"></stop><stop offset="1" stop-color="#A9A9AA"></stop></linearGradient><linearGradient id="star-g3" x1="12.7076" y1="0" x2="12.7076" y2="40.0895" gradientUnits="userSpaceOnUse"><stop stop-color="white"></stop><stop offset="1" stop-color="#E6E6E7"></stop></linearGradient><linearGradient id="star-g4" x1="26.8553" y1="0" x2="26.8553" y2="25.4153" gradientUnits="userSpaceOnUse"><stop stop-color="#F3F3F3"></stop><stop offset="1" stop-color="#A9A9AA"></stop></linearGradient></defs></svg><svg class="h-[14px] w-auto flex-shrink-0" viewBox="54 12 139 23" fill="none" xmlns="http://www.w3.org/2000/svg" aria-label="FutureAGI"><path d="M54.7168 34.0518V12.8484H68.1504V15.4099H57.506V22.2689H67.1542V24.8304H57.506V34.0518H54.7168Z" fill="white"></path><path d="M76.1548 34.3933C75.0543 34.3933 74.0582 34.1372 73.1664 33.6249C72.2936 33.1126 71.6105 32.4011 71.1172 31.4903C70.6429 30.5606 70.4057 29.498 70.4057 28.3027V18.7113H73.0525V28.0181C73.0525 28.777 73.2043 29.4411 73.5079 30.0103C73.8305 30.5795 74.2669 31.0254 74.8171 31.348C75.3863 31.6706 76.0315 31.8318 76.7525 31.8318C77.4735 31.8318 78.1091 31.6706 78.6594 31.348C79.2286 31.0254 79.665 30.5606 79.9686 29.9534C80.2911 29.3462 80.4524 28.6252 80.4524 27.7904V18.7113H83.1277V34.0518H80.5378V31.0634L80.9647 31.3195C80.6042 32.2872 79.9875 33.0462 79.1147 33.5964C78.2609 34.1277 77.2743 34.3933 76.1548 34.3933Z" fill="white"></path><path d="M93.4613 34.2226C91.9623 34.2226 90.8049 33.7956 89.989 32.9418C89.1921 32.088 88.7937 30.8831 88.7937 29.3273V21.2444H86.0045V18.7113H86.5737C87.2568 18.7113 87.7975 18.5026 88.196 18.0852C88.5945 17.6678 88.7937 17.1175 88.7937 16.4344V15.1822H91.4406V18.7113H94.8843V21.2444H91.4406V29.2419C91.4406 29.7542 91.5164 30.2001 91.6682 30.5795C91.839 30.959 92.1141 31.2626 92.4936 31.4903C92.8731 31.699 93.3759 31.8034 94.002 31.8034C94.1349 31.8034 94.2961 31.7939 94.4859 31.7749C94.6946 31.7559 94.8843 31.737 95.0551 31.718V34.0518C94.8084 34.1087 94.5333 34.1467 94.2297 34.1656C93.9261 34.2036 93.67 34.2226 93.4613 34.2226Z" fill="white"></path><path d="M103.949 34.3933C102.848 34.3933 101.852 34.1372 100.96 33.6249C100.088 33.1126 99.4044 32.4011 98.9111 31.4903C98.4368 30.5606 98.1996 29.498 98.1996 28.3027V18.7113H100.846V28.0181C100.846 28.777 100.998 29.4411 101.302 30.0103C101.624 30.5795 102.061 31.0254 102.611 31.348C103.18 31.6706 103.825 31.8318 104.546 31.8318C105.267 31.8318 105.903 31.6706 106.453 31.348C107.022 31.0254 107.459 30.5606 107.762 29.9534C108.085 29.3462 108.246 28.6252 108.246 27.7904V18.7113H110.922V34.0518H108.332V31.0634L108.759 31.3195C108.398 32.2872 107.781 33.0462 106.909 33.5964C106.055 34.1277 105.068 34.3933 103.949 34.3933Z" fill="white"></path><path d="M115.022 34.0518V18.7113H117.612V21.529L117.328 21.1305C117.688 20.2577 118.238 19.6126 118.978 19.1952C119.718 18.7588 120.62 18.5406 121.682 18.5406H122.621V21.0451H121.284C120.202 21.0451 119.329 21.3867 118.665 22.0697C118.001 22.7338 117.669 23.6825 117.669 24.9158V34.0518H115.022Z" fill="white"></path><path d="M132.031 34.3933C130.551 34.3933 129.232 34.0423 128.075 33.3403C126.917 32.6382 126.007 31.68 125.342 30.4657C124.678 29.2324 124.346 27.8568 124.346 26.3389C124.346 24.802 124.669 23.4358 125.314 22.2405C125.978 21.0451 126.87 20.1059 127.989 19.4229C129.128 18.7208 130.399 18.3698 131.803 18.3698C132.942 18.3698 133.947 18.5785 134.82 18.9959C135.712 19.3944 136.461 19.9446 137.068 20.6467C137.695 21.3297 138.169 22.1172 138.491 23.0089C138.833 23.8817 139.004 24.7925 139.004 25.7412C139.004 25.9499 138.985 26.1871 138.947 26.4527C138.928 26.6994 138.899 26.9365 138.861 27.1642H126.282V24.8874H137.325L136.072 25.912C136.243 24.9253 136.148 24.043 135.788 23.2651C135.427 22.4871 134.896 21.8705 134.194 21.4151C133.492 20.9597 132.695 20.7321 131.803 20.7321C130.911 20.7321 130.095 20.9597 129.355 21.4151C128.615 21.8705 128.037 22.5251 127.619 23.3789C127.221 24.2138 127.06 25.2099 127.135 26.3673C127.06 27.4868 127.23 28.4734 127.648 29.3273C128.084 30.1621 128.691 30.8167 129.469 31.2911C130.266 31.7464 131.13 31.9741 132.059 31.9741C133.084 31.9741 133.947 31.737 134.649 31.2626C135.351 30.7883 135.92 30.1811 136.357 29.4411L138.577 30.5795C138.273 31.2816 137.799 31.9267 137.154 32.5149C136.528 33.0841 135.778 33.5395 134.905 33.881C134.052 34.2226 133.093 34.3933 132.031 34.3933Z" fill="white"></path><path d="M145.638 34.0518L153.237 12.8484H156.539L164.138 34.0518H161.149L159.413 29.0711H150.363L148.627 34.0518H145.638ZM151.245 26.5096H158.531L154.49 14.8691H155.286L151.245 26.5096Z" fill="white"></path><path d="M175.739 34.3933C174.24 34.3933 172.855 34.1277 171.583 33.5964C170.312 33.0462 169.212 32.2777 168.282 31.2911C167.352 30.3044 166.622 29.147 166.09 27.8188C165.578 26.4907 165.322 25.0391 165.322 23.4643C165.322 21.8705 165.578 20.4095 166.09 19.0813C166.603 17.7531 167.324 16.5957 168.253 15.6091C169.183 14.6224 170.284 13.8635 171.555 13.3322C172.826 12.782 174.211 12.5068 175.71 12.5068C177.171 12.5068 178.48 12.763 179.638 13.2753C180.814 13.7876 181.801 14.4517 182.598 15.2675C183.414 16.0834 183.992 16.9562 184.334 17.886L181.829 19.1098C181.336 17.8765 180.567 16.8993 179.524 16.1783C178.48 15.4573 177.209 15.0968 175.71 15.0968C174.23 15.0968 172.911 15.4478 171.754 16.1498C170.616 16.8519 169.724 17.829 169.079 19.0813C168.434 20.3336 168.111 21.7946 168.111 23.4643C168.111 25.115 168.434 26.5666 169.079 27.8188C169.743 29.0711 170.644 30.0483 171.783 30.7503C172.94 31.4524 174.259 31.8034 175.739 31.8034C177.029 31.8034 178.196 31.5282 179.239 30.978C180.283 30.4278 181.108 29.6688 181.715 28.7011C182.323 27.7335 182.626 26.614 182.626 25.3427V24.0335L183.907 25.2289H175.71V22.8097H185.444V24.6881C185.444 26.1681 185.188 27.5058 184.675 28.7011C184.163 29.8965 183.461 30.9211 182.569 31.7749C181.677 32.6098 180.643 33.2549 179.467 33.7103C178.291 34.1656 177.048 34.3933 175.739 34.3933Z" fill="white"></path><path d="M189.212 34.0518V12.8484H192.001V34.0518H189.212Z" fill="white"></path></svg></div> <span class="text-[13px] text-[#a1a1aa]" data-astro-cid-35ed7um5>
&copy; 2026 Future AGI, Inc.
</span> </div> <!-- Legal Links --> <div class="flex flex-wrap items-center gap-x-6 gap-y-2" data-astro-cid-35ed7um5> <a href="/terms/" class="text-[13px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Terms of use </a><a href="/privacy/" class="text-[13px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Privacy policy </a><a href="/security/" class="text-[13px] text-[#a1a1aa] hover:text-[#fafafa] transition-colors" data-astro-cid-35ed7um5> Security </a> </div> </div> </div> </div> </footer>   <script>
  (function () {
    const FORMATTERS = {
      compact: (n) => {
        if (typeof n !== 'number' || !isFinite(n)) return null;
        if (n < 1000) return String(n);
        if (n < 10000) return (n / 1000).toFixed(1).replace(/\.0$/, '') + 'k';
        if (n < 1_000_000) return Math.round(n / 1000) + 'k';
        return (n / 1_000_000).toFixed(1).replace(/\.0$/, '') + 'm';
      },
      full: (n) => {
        if (typeof n !== 'number' || !isFinite(n)) return null;
        return n.toLocaleString('en-US');
      },
    };

    function applyStats(stats) {
      const els = document.querySelectorAll('[data-stat]');
      if (!els.length) return;

      els.forEach((el) => {
        const key = el.getAttribute('data-stat');
        const fmt = el.getAttribute('data-stat-format') || 'compact';
        const raw = stats[key];
        if (raw === undefined || raw === null) return;

        const formatter = FORMATTERS[fmt] || FORMATTERS.compact;
        const next = formatter(raw);
        if (next === null) return;

        const current = (el.textContent || '').trim();
        if (current === next) return;

        // Soft fade. If reduced-motion, just swap.
        if (window.matchMedia('(prefers-reduced-motion: reduce)').matches) {
          el.textContent = next;
          return;
        }
        el.style.transition = 'opacity 0.25s ease-out';
        el.style.opacity = '0.4';
        setTimeout(() => {
          el.textContent = next;
          el.style.opacity = '1';
          // Trigger the amber pulse on any element that opted in
          // (currently the header star-pill count).
          if (el.classList.contains('star-pill-count')) {
            el.classList.remove('is-updated');
            // Force a reflow so the animation re-fires.
            void el.offsetWidth;
            el.classList.add('is-updated');
          }
        }, 200);
      });

      // Expose for debugging — surfaces on `window.__fagiStats` in DevTools.
      try { window.__fagiStats = stats; } catch (_) {}
    }

    async function hydrate() {
      try {
        const res = await fetch('/stars.json', { cache: 'no-cache' });
        if (!res.ok) return;
        const data = await res.json();
        if (typeof data !== 'object' || !data) return;
        applyStats(data);
      } catch (_) {
        // Network/parse error — leave the build-time values alone.
      }
    }

    if (document.readyState === 'loading') {
      document.addEventListener('DOMContentLoaded', hydrate, { once: true });
    } else {
      hydrate();
    }
    // Astro view transitions re-fire this event when navigating.
    document.addEventListener('astro:page-load', hydrate);
  })();
</script> <!-- Scroll Navigation Arrows --> <div id="scroll-nav" class="scroll-nav" data-astro-cid-sckkx6r4> <button id="scroll-top" class="scroll-nav-btn" title="Scroll to top" aria-label="Scroll to top" data-astro-cid-sckkx6r4> <svg width="14" height="14" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-sckkx6r4><path d="M18 15l-6-6-6 6" data-astro-cid-sckkx6r4></path></svg> </button> <button id="scroll-bottom" class="scroll-nav-btn" title="Scroll to bottom" aria-label="Scroll to bottom" data-astro-cid-sckkx6r4> <svg width="14" height="14" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2" stroke-linecap="round" stroke-linejoin="round" data-astro-cid-sckkx6r4><path d="M6 9l6 6 6-6" data-astro-cid-sckkx6r4></path></svg> </button> </div>  <script type="module" src="/_astro/Layout.astro_astro_type_script_index_0_lang.Dm6q9DDa.js"></script> <!-- AI Chat Widget + Floating Button + Question Bubbles --> <!-- Design token bridge: product-docs → landing page --><!-- Chat dependencies loaded lazily when widget opens (saves ~350KB on initial page load) --><script>
  window._chatDepsLoaded = false;
  window._chatDepsLoading = false;
  window.loadChatDeps = function() {
    if (window._chatDepsLoaded || window._chatDepsLoading) return Promise.resolve();
    window._chatDepsLoading = true;
    var loads = [];
    // marked.js
    if (typeof marked === 'undefined') {
      loads.push(new Promise(function(r) { var s = document.createElement('script'); s.src = 'https://cdn.jsdelivr.net/npm/marked/marked.min.js'; s.onload = r; s.onerror = r; document.head.appendChild(s); }));
    }
    // highlight.js
    if (typeof hljs === 'undefined') {
      loads.push(new Promise(function(r) { var s = document.createElement('script'); s.src = 'https://cdn.jsdelivr.net/gh/highlightjs/cdn-release@11.9.0/build/highlight.min.js'; s.onload = r; s.onerror = r; document.head.appendChild(s); }));
      var link = document.createElement('link'); link.rel = 'stylesheet'; link.href = 'https://cdn.jsdelivr.net/gh/highlightjs/cdn-release@11.9.0/build/styles/github-dark.min.css'; document.head.appendChild(link);
    }
    // Cloudflare Turnstile
    if (typeof turnstile === 'undefined') {
      var ts = document.createElement('script'); ts.src = 'https://challenges.cloudflare.com/turnstile/v0/api.js?render=explicit'; ts.async = true; document.head.appendChild(ts);
    }
    return Promise.all(loads).then(function() { window._chatDepsLoaded = true; window._chatDepsLoading = false; });
  };
</script><!-- Chat popup --><div id="ai-chat-popup" class="fixed z-50 rounded-2xl border border-[var(--color-border-default)] shadow-2xl shadow-black/30 flex-col overflow-hidden hidden" style="background: #111111; bottom: 24px; right: 24px; width: 420px; height: 560px;" data-api-url="https://docs-api.futureagi.com" data-turnstile-key="0x4AAAAAACts5xVgAJ5y5Vm_"> <!-- Resize handles: corner + edges --> <div id="ai-resize-handle" class="absolute top-0 left-0 w-4 h-4 z-10" style="cursor: nw-resize;"> <svg class="w-3 h-3 text-[var(--color-text-muted)] opacity-0 hover:opacity-100 transition-opacity m-0.5 rotate-90" viewBox="0 0 24 24" fill="currentColor"> <circle cx="4" cy="4" r="2"></circle><circle cx="12" cy="4" r="2"></circle><circle cx="4" cy="12" r="2"></circle> </svg> </div> <div id="ai-resize-top" class="absolute top-0 left-4 right-0 h-2 z-10" style="cursor: n-resize;"></div> <div id="ai-resize-left" class="absolute top-4 left-0 bottom-0 w-2 z-10" style="cursor: w-resize;"></div> <!-- Header --> <div class="flex items-center gap-2.5 px-4 py-2.5 border-b border-[var(--color-border-subtle)] flex-shrink-0" style="background: #161616;"> <div class="w-5 h-5 rounded-md flex items-center justify-center flex-shrink-0" style="background: #8b5cf6;"> <svg class="w-3 h-3" viewBox="0 0 24 24" fill="white" stroke="none"> <path d="M9.663 17h4.673M12 3v1m6.364 1.636l-.707.707M21 12h-1M4 12H3m3.343-5.657l-.707-.707m2.828 9.9a5 5 0 117.072 0l-.548.547A3.374 3.374 0 0014 18.469V19a2 2 0 11-4 0v-.531c0-.895-.356-1.754-.988-2.386l-.548-.547z"></path> </svg> </div> <span class="text-sm font-semibold text-[var(--color-text-primary)] flex-1">AI Assistant</span> <span class="px-1.5 py-0.5 text-[10px] font-medium rounded bg-purple-500/15 text-purple-400">Beta</span> <button id="ai-new-chat-btn" class="p-1 rounded-md text-[var(--color-text-muted)] hover:text-[var(--color-text-primary)] hover:bg-[var(--color-bg-hover)] transition-colors cursor-pointer" title="New chat"> <svg class="w-3.5 h-3.5" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2"><path d="M12 5v14M5 12h14" stroke-linecap="round" stroke-linejoin="round"></path></svg> </button> <button id="ai-sidebar-btn" class="p-1 rounded-md text-[var(--color-text-muted)] hover:text-[var(--color-text-primary)] hover:bg-[var(--color-bg-hover)] transition-colors cursor-pointer" title="Dock to sidebar"> <!-- Sidebar icon (shown in popup mode) --> <svg id="ai-sb-dock" class="w-3.5 h-3.5" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2"><path d="M9 3h12v18H9M3 3h6v18H3" stroke-linecap="round" stroke-linejoin="round"></path></svg> <!-- Floating popup icon (shown in sidebar mode) --> <svg id="ai-sb-undock" class="w-3.5 h-3.5 hidden" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2"><rect x="4" y="4" width="12" height="12" rx="2" stroke-linecap="round" stroke-linejoin="round"></rect><path d="M20 8v12a2 2 0 01-2 2H8" stroke-linecap="round" stroke-linejoin="round"></path></svg> </button> <button id="ai-fullscreen-btn" class="p-1 rounded-md text-[var(--color-text-muted)] hover:text-[var(--color-text-primary)] hover:bg-[var(--color-bg-hover)] transition-colors cursor-pointer" title="Toggle fullscreen"> <svg id="ai-fs-expand" class="w-3.5 h-3.5" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2"><path d="M15 3h6v6M9 21H3v-6M21 3l-7 7M3 21l7-7" stroke-linecap="round" stroke-linejoin="round"></path></svg> <svg id="ai-fs-shrink" class="w-3.5 h-3.5 hidden" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2"><path d="M4 14h6v6M20 10h-6V4M14 10l7-7M3 21l7-7" stroke-linecap="round" stroke-linejoin="round"></path></svg> </button> <button id="ai-popup-close" class="p-1 rounded-md text-[var(--color-text-muted)] hover:text-[var(--color-text-primary)] hover:bg-[var(--color-bg-hover)] transition-colors cursor-pointer" title="Close"> <svg class="w-3.5 h-3.5" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2"><path d="M6 18L18 6M6 6l12 12" stroke-linecap="round" stroke-linejoin="round"></path></svg> </button> </div> <!-- Messages --> <div id="ai-widget-messages" class="flex-1 overflow-y-auto px-4 py-4 hide-scrollbar"> <!-- Empty state --> <div id="ai-empty-state" class="flex flex-col items-center justify-center h-full text-center px-4"> <div class="w-10 h-10 rounded-xl flex items-center justify-center mb-3" style="background: #8b5cf6;"> <svg class="w-5 h-5" viewBox="0 0 24 24" fill="white" stroke="none"> <path d="M9.663 17h4.673M12 3v1m6.364 1.636l-.707.707M21 12h-1M4 12H3m3.343-5.657l-.707-.707m2.828 9.9a5 5 0 117.072 0l-.548.547A3.374 3.374 0 0014 18.469V19a2 2 0 11-4 0v-.531c0-.895-.356-1.754-.988-2.386l-.548-.547z"></path> </svg> </div> <div class="text-sm font-semibold text-[var(--color-text-primary)] mb-1">FutureAGI AI Assistant</div> <p class="text-xs text-[var(--color-text-muted)] mb-4 leading-relaxed max-w-[240px]">
Ask me anything about the FutureAGI platform — I can search across all docs instantly.
</p> <div id="ai-widget-quick" class="space-y-1.5 w-full max-w-[280px]"> <button class="ai-wq w-full text-left px-3 py-2.5 text-[13px] border rounded-lg text-[var(--color-text-secondary)] hover:text-[var(--color-text-primary)] transition-all cursor-pointer" style="background: rgba(139,92,246,0.08); border-color: rgba(139,92,246,0.25);" data-q="What is FutureAGI and how does it help with AI agent reliability?"><svg class="inline-block w-3.5 h-3.5 mr-1.5 -mt-0.5" viewBox="0 0 24 24" fill="none" stroke="currentColor" stroke-width="2"><path d="M9.663 17h4.673M12 3v1m6.364 1.636l-.707.707M21 12h-1M4 12H3m3.343-5.657l-.707-.707m2.828 9.9a5 5 0 117.072 0l-.548.547A3.374 3.374 0 0014 18.469V19a2 2 0 11-4 0v-.531c0-.895-.356-1.754-.988-2.386l-.548-.547z"></path></svg>What is FutureAGI?</button> <button class="ai-wq w-full text-left px-3 py-2.5 text-[13px] border border-[var(--color-border-subtle)] rounded-lg text-[var(--color-text-secondary)] hover:border-[var(--color-border-default)] hover:text-[var(--color-text-primary)] transition-all cursor-pointer" style="background: #161616;" data-q="What can FutureAGI do? Give me an overview of all features.">What can FutureAGI do?</button> <button class="ai-wq w-full text-left px-3 py-2.5 text-[13px] border border-[var(--color-border-subtle)] rounded-lg text-[var(--color-text-secondary)] hover:border-[var(--color-border-default)] hover:text-[var(--color-text-primary)] transition-all cursor-pointer" style="background: #161616;" data-q="How do I run my first evaluation?">How do I run my first evaluation?</button> <button class="ai-wq w-full text-left px-3 py-2.5 text-[13px] border border-[var(--color-border-subtle)] rounded-lg text-[var(--color-text-secondary)] hover:border-[var(--color-border-default)] hover:text-[var(--color-text-primary)] transition-all cursor-pointer" style="background: #161616;" data-q="How do I set up tracing for my AI app?">How do I set up tracing?</button> <button class="ai-wq w-full text-left px-3 py-2.5 text-[13px] border border-[var(--color-border-subtle)] rounded-lg text-[var(--color-text-secondary)] hover:border-[var(--color-border-default)] hover:text-[var(--color-text-primary)] transition-all cursor-pointer" style="background: #161616;" data-q="How do I detect hallucinations in my RAG pipeline?">How do I detect hallucinations?</button> </div> </div> </div> <!-- Stop button (deprecated — send button toggles to stop icon during generation) --> <div id="ai-stop-wrap" class="hidden" style="display:none !important;"></div> <!-- Input --> <div class="px-4 py-3 border-t border-[var(--color-border-subtle)] flex-shrink-0" style="background: #161616;"> <div id="ai-input-wrap" class="flex items-end gap-2 border border-[var(--color-border-subtle)] rounded-xl px-3 py-2.5 focus-within:border-[#8b5cf6] transition-colors" style="background: #111111;"> <textarea id="ai-widget-input" placeholder="Ask a question..." rows="1" maxlength="1000" class="flex-1 bg-transparent border-none resize-none text-[13px] text-[var(--color-text-primary)] placeholder:text-[var(--color-text-muted)] leading-snug" style="min-height: 20px; max-height: 80px; outline: none !important; box-shadow: none !important;"></textarea> <button id="ai-widget-send" class="w-7 h-7 flex items-center justify-center rounded-lg flex-shrink-0 cursor-pointer hover:opacity-90 transition-opacity" style="background: #8b5cf6;"> <svg class="w-3.5 h-3.5" viewBox="0 0 24 24" fill="white" stroke="none"> <path d="M2.01 21L23 12 2.01 3 2 10l15 2-15 2z"></path> </svg> </button> </div> <p class="text-[10px] text-[var(--color-text-muted)] text-center mt-2">Built by FAGI with ❤️</p> </div> </div> <script type="module" src="/_astro/AiChatWidget.astro_astro_type_script_index_0_lang.BK0oVwqu.js"></script>  <!-- Global styles for citation highlighting (needs to target elements outside the widget) -->  <!-- Floating question bubbles (appear one at a time above the FAB) --> <div id="ai-fab-bubbles" class="ai-fab-bubbles" data-astro-cid-sckkx6r4></div> <button id="ai-chat-fab" onclick="window.openAiChat && window.openAiChat()" class="ai-fab-btn" title="Ask AI" data-astro-cid-sckkx6r4> <!-- Outer glow ring --> <div class="ai-fab-glow" data-astro-cid-sckkx6r4></div> <!-- Hand-drawn constellation "?" --> <svg class="ai-fab-svg" viewBox="0 0 44 44" fill="none" data-astro-cid-sckkx6r4> <!-- Ambient twinkling stars --> <circle cx="4" cy="6" r="0.7" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.15;0.7;0.15" dur="3s" repeatCount="indefinite" data-astro-cid-sckkx6r4></animate></circle> <circle cx="40" cy="4" r="0.5" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.1;0.55;0.1" dur="2.6s" repeatCount="indefinite" begin="0.4s" data-astro-cid-sckkx6r4></animate></circle> <circle cx="3" cy="37" r="0.6" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.1;0.6;0.1" dur="3.8s" repeatCount="indefinite" begin="1.2s" data-astro-cid-sckkx6r4></animate></circle> <circle cx="41" cy="39" r="0.5" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.08;0.5;0.08" dur="3.2s" repeatCount="indefinite" begin="1.8s" data-astro-cid-sckkx6r4></animate></circle> <circle cx="38" cy="21" r="0.4" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.1;0.5;0.1" dur="2.9s" repeatCount="indefinite" begin="0.7s" data-astro-cid-sckkx6r4></animate></circle> <circle cx="6" cy="27" r="0.5" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.12;0.5;0.12" dur="3.4s" repeatCount="indefinite" begin="2.1s" data-astro-cid-sckkx6r4></animate></circle> <circle cx="35" cy="30" r="0.4" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.08;0.45;0.08" dur="4s" repeatCount="indefinite" begin="0.3s" data-astro-cid-sckkx6r4></animate></circle> <circle cx="9" cy="16" r="0.35" fill="white" data-astro-cid-sckkx6r4><animate attributeName="opacity" values="0.1;0.4;0.1" dur="3.1s" repeatCount="indefinite" begin="1.5s" data-astro-cid-sckkx6r4></animate></circle> <!-- Constellation path --> <path d="M13.5 14.5 C14 11, 16 8.5, 19 7.5 C22 6.5, 26 7, 28 9.5 C30 12, 29.5 15, 27.5 17.5 C25.5 20, 23 22, 22 25" stroke="url(#fab-line-grad)" stroke-width="0.7" stroke-linecap="round" fill="none" opacity="0.6" stroke-dasharray="2 3" data-astro-cid-sckkx6r4></path> <path d="M14 13.5 C15 10, 17.5 8, 20 7 C24 6, 27.5 8, 28.5 11 C29.5 14, 28 17, 26 19.5 C24 22, 22.5 23.5, 22 25.5" stroke="rgba(167,139,250,0.2)" stroke-width="0.4" stroke-linecap="round" fill="none" data-astro-cid-sckkx6r4></path> <!-- Star nodes --> <circle cx="13.5" cy="14.5" r="1.5" fill="white" opacity="0.8" data-astro-cid-sckkx6r4></circle> <line x1="12" y1="14.5" x2="15" y2="14.5" stroke="white" stroke-width="0.25" opacity="0.3" data-astro-cid-sckkx6r4></line> <line x1="13.5" y1="13" x2="13.5" y2="16" stroke="white" stroke-width="0.25" opacity="0.3" data-astro-cid-sckkx6r4></line> <circle cx="16" cy="9.5" r="1.2" fill="white" opacity="0.65" data-astro-cid-sckkx6r4></circle> <circle cx="20" cy="7" r="2" fill="white" data-astro-cid-sckkx6r4></circle> <line x1="17.5" y1="7" x2="22.5" y2="7" stroke="white" stroke-width="0.35" opacity="0.45" data-astro-cid-sckkx6r4></line> <line x1="20" y1="4.5" x2="20" y2="9.5" stroke="white" stroke-width="0.35" opacity="0.45" data-astro-cid-sckkx6r4></line> <line x1="18.2" y1="5.2" x2="21.8" y2="8.8" stroke="white" stroke-width="0.2" opacity="0.2" data-astro-cid-sckkx6r4></line> <line x1="21.8" y1="5.2" x2="18.2" y2="8.8" stroke="white" stroke-width="0.2" opacity="0.2" data-astro-cid-sckkx6r4></line> <circle cx="25.5" cy="8" r="1.3" fill="white" opacity="0.7" data-astro-cid-sckkx6r4></circle> <circle cx="28.5" cy="11" r="1.5" fill="white" opacity="0.8" data-astro-cid-sckkx6r4></circle> <line x1="27" y1="11" x2="30" y2="11" stroke="white" stroke-width="0.25" opacity="0.3" data-astro-cid-sckkx6r4></line> <line x1="28.5" y1="9.5" x2="28.5" y2="12.5" stroke="white" stroke-width="0.25" opacity="0.3" data-astro-cid-sckkx6r4></line> <circle cx="28" cy="16" r="1.1" fill="white" opacity="0.6" data-astro-cid-sckkx6r4></circle> <circle cx="25.5" cy="20" r="1.3" fill="white" opacity="0.7" data-astro-cid-sckkx6r4></circle> <circle cx="22" cy="25" r="1.5" fill="white" opacity="0.8" data-astro-cid-sckkx6r4></circle> <line x1="20.5" y1="25" x2="23.5" y2="25" stroke="white" stroke-width="0.25" opacity="0.3" data-astro-cid-sckkx6r4></line> <line x1="22" y1="23.5" x2="22" y2="26.5" stroke="white" stroke-width="0.25" opacity="0.3" data-astro-cid-sckkx6r4></line> <!-- The DOT - pulsing purple orb (more glow) --> <circle cx="22" cy="34" r="7" fill="#8b5cf6" opacity="0.06" data-astro-cid-sckkx6r4> <animate attributeName="r" values="5;9;5" dur="2.5s" repeatCount="indefinite" data-astro-cid-sckkx6r4></animate> <animate attributeName="opacity" values="0.06;0.02;0.06" dur="2.5s" repeatCount="indefinite" data-astro-cid-sckkx6r4></animate> </circle> <circle cx="22" cy="34" r="4.5" fill="#8b5cf6" opacity="0.12" data-astro-cid-sckkx6r4> <animate attributeName="opacity" values="0.12;0.05;0.12" dur="2s" repeatCount="indefinite" data-astro-cid-sckkx6r4></animate> </circle> <circle cx="22" cy="34" r="2.8" fill="#a78bfa" data-astro-cid-sckkx6r4></circle> <circle cx="22" cy="34" r="1.2" fill="white" opacity="0.35" data-astro-cid-sckkx6r4></circle> <circle cx="21" cy="33" r="0.5" fill="white" opacity="0.6" data-astro-cid-sckkx6r4></circle> <defs data-astro-cid-sckkx6r4> <linearGradient id="fab-line-grad" x1="13" y1="14" x2="22" y2="25" data-astro-cid-sckkx6r4> <stop offset="0%" stop-color="#8b5cf6" stop-opacity="0.6" data-astro-cid-sckkx6r4></stop> <stop offset="100%" stop-color="#a78bfa" stop-opacity="0.35" data-astro-cid-sckkx6r4></stop> </linearGradient> </defs> </svg> </button>  <script>
      (function() {
        // 3 generic questions (always available)
        var genericQuestions = [
          "What can FutureAGI do?",
          "How do I detect hallucinations?",
          "How do I run my first evaluation?"
        ];

        // Build 2 page-context questions from current page content
        function getPageQuestions() {
          var pageQs = [];
          var path = window.location.pathname.replace(/\/$/, '') || '/';

          // Try to read the visible section headings
          var h1 = document.querySelector('h1');
          var pageTitle = h1 ? h1.textContent.trim() : '';

          // Find the section the user is currently looking at
          var visibleHeading = '';
          var h2s = document.querySelectorAll('h2');
          for (var i = 0; i < h2s.length; i++) {
            var rect = h2s[i].getBoundingClientRect();
            if (rect.top > 0 && rect.top < window.innerHeight * 0.6) {
              visibleHeading = h2s[i].textContent.trim();
              break;
            }
          }
          // Fallback: use last h2 above viewport center
          if (!visibleHeading) {
            for (var j = h2s.length - 1; j >= 0; j--) {
              if (h2s[j].getBoundingClientRect().top < window.innerHeight * 0.5) {
                visibleHeading = h2s[j].textContent.trim();
                break;
              }
            }
          }

          // Generate contextual questions
          if (visibleHeading) {
            pageQs.push('Tell me about "' + truncate(visibleHeading, 28) + '"');
          }
          if (pageTitle && pageTitle !== visibleHeading) {
            pageQs.push('Explain "' + truncate(pageTitle, 28) + '"');
          }

          // If we couldn't get enough from headings, use path-based hints
          if (pageQs.length < 2) {
            var segments = path.split('/').filter(Boolean);
            if (segments.length > 0) {
              var topic = segments[segments.length - 1].replace(/-/g, ' ');
              topic = topic.charAt(0).toUpperCase() + topic.slice(1);
              if (topic.length > 2 && topic.toLowerCase() !== 'home') {
                pageQs.push('How does ' + truncate(topic, 28) + ' work?');
              }
            }
          }

          // Fallback page questions if nothing detected
          while (pageQs.length < 2) {
            var fallbacks = [
              "How do I set up tracing?",
              "What evaluators are available?",
              "How does guardrails work?"
            ];
            pageQs.push(fallbacks[pageQs.length]);
          }

          return pageQs.slice(0, 2);
        }

        function truncate(str, max) {
          return str.length > max ? str.substring(0, max - 1) + '…' : str;
        }

        // Build the final question list: 2 page-context first, then 3 generic
        function buildQuestions() {
          var pageQs = getPageQuestions();
          return pageQs.concat(genericQuestions);
        }

        var bubbleContainer, fab, popup, currentBubble = null;
        var bubbleTimer = null, cycleTimer = null;
        var qIndex = 0;
        var questions = [];
        var dismissed = false;

        function showBubble() {
          if (dismissed) return;
          if (!bubbleContainer || !fab) return;
          if (fab.style.display === 'none') return;

          // Build questions on first bubble (allows page to fully render)
          if (questions.length === 0) questions = buildQuestions();

          // Remove old bubble
          if (currentBubble) {
            currentBubble.style.animation = 'aiBubbleOut 0.3s ease-in forwards';
            var old = currentBubble;
            setTimeout(function() { if (old.parentNode) old.parentNode.removeChild(old); }, 300);
          }

          var bubble = document.createElement('div');
          bubble.className = 'ai-fab-bubble';
          bubble.textContent = questions[qIndex];
          bubble.setAttribute('data-q', questions[qIndex]);
          bubble.addEventListener('click', function() {
            var q = this.getAttribute('data-q');
            dismissed = true;
            clearTimeout(bubbleTimer);
            clearTimeout(cycleTimer);
            hideBubbles();
            if (window.openAiChat) window.openAiChat();
            setTimeout(function() {
              var input = document.getElementById('ai-widget-input');
              if (input) {
                input.value = q;
                var sendBtn = document.getElementById('ai-widget-send');
                if (sendBtn) sendBtn.click();
              }
            }, 300);
          });

          bubbleContainer.appendChild(bubble);
          currentBubble = bubble;
          qIndex = (qIndex + 1) % questions.length;

          bubbleTimer = setTimeout(function() {
            if (currentBubble) {
              currentBubble.style.animation = 'aiBubbleOut 0.3s ease-in forwards';
              var old = currentBubble;
              setTimeout(function() { if (old.parentNode) old.parentNode.removeChild(old); }, 300);
              currentBubble = null;
            }
            cycleTimer = setTimeout(showBubble, 3000);
          }, 5000);
        }

        function hideBubbles() {
          if (bubbleContainer) bubbleContainer.innerHTML = '';
          currentBubble = null;
        }

        function syncFab() {
          if (!fab || !popup) return;
          var isVisible = popup.classList.contains('flex');
          fab.style.display = isVisible ? 'none' : 'flex';
          bubbleContainer.style.display = isVisible ? 'none' : 'flex';
          if (isVisible) {
            hideBubbles();
            clearTimeout(bubbleTimer);
            clearTimeout(cycleTimer);
          }
        }

        function init() {
          fab = document.getElementById('ai-chat-fab');
          popup = document.getElementById('ai-chat-popup');
          bubbleContainer = document.getElementById('ai-fab-bubbles');
          if (!fab || !popup) return;

          // Reset questions on page change (new page context)
          questions = [];
          qIndex = 0;

          var observer = new MutationObserver(syncFab);
          observer.observe(popup, { attributes: true, attributeFilter: ['class'] });
          syncFab();

          if (!dismissed) {
            clearTimeout(bubbleTimer);
            clearTimeout(cycleTimer);
            setTimeout(showBubble, 2000);
          }
        }

        init();
        document.addEventListener('astro:page-load', init);

        // Rebuild questions when user scrolls to a new section (debounced)
        var scrollDebounce = null;
        window.addEventListener('scroll', function() {
          clearTimeout(scrollDebounce);
          scrollDebounce = setTimeout(function() {
            if (!dismissed && questions.length > 0) {
              // Refresh the page-context questions (first 2 slots)
              var fresh = getPageQuestions();
              questions[0] = fresh[0];
              questions[1] = fresh[1];
            }
          }, 500);
        }, { passive: true });
      })();
    </script> <!-- Initialize smooth scroll and animations --> <script type="module" src="/_astro/Layout.astro_astro_type_script_index_1_lang.DDcg2uOl.js"></script> <script type="module">function l(){const t=document.getElementById("scroll-nav"),o=document.getElementById("scroll-top"),e=document.getElementById("scroll-bottom");if(!t||!o||!e)return;const n=()=>{window.scrollY>300?t.classList.add("visible"):t.classList.remove("visible")};window.addEventListener("scroll",n,{passive:!0}),n(),o.addEventListener("click",()=>{window.scrollTo({top:0,behavior:"smooth"})}),e.addEventListener("click",()=>{window.scrollTo({top:document.body.scrollHeight,behavior:"smooth"})})}document.addEventListener("DOMContentLoaded",l);document.addEventListener("astro:page-load",l);</script> </body> </html>  <script type="module" src="/_astro/startups.astro_astro_type_script_index_0_lang.DE2Fzk6z.js"></script> <script type="module" src="/_astro/startups.astro_astro_type_script_index_1_lang.DcLI9KJB.js"></script> <script type="module">function b(){const f=document.getElementById("rocket-canvas");if(!f)return;const a=f.getContext("2d");if(!a)return;const u=Math.min(window.devicePixelRatio||1,2);let p,x,d=0,C;const M=10,c=16,m="0123456789ABCDEFabcdef.:+=-~|",y="|!1lI/\\";let s,h,w=[],g=[];const T=60;function E(){p=f.clientWidth,x=f.clientHeight,f.width=p*u,f.height=x*u,a.setTransform(u,0,0,u,0,0),s=Math.ceil(p/M),h=Math.ceil(x/c),w=[];for(let e=0;e<h;e++){const t=[];for(let n=0;n<s;n++)t.push({char:m[Math.floor(Math.random()*m.length)],alpha:0,timer:Math.floor(Math.random()*200),interval:80+Math.floor(Math.random()*250),speed:.5+Math.random()*2});w.push(t)}}function B(){if(g.length>=T)return;const e=s/2,t=s*.18,n=e+(Math.random()-.5)*t*2,o=h+Math.random()*5,l=80+Math.random()*160;g.push({x:n,y:o,speed:.15+Math.random()*.35,life:0,maxLife:l,char:y[Math.floor(Math.random()*y.length)],brightness:.4+Math.random()*.6,drift:(Math.random()-.5)*.03})}function F(e,t){const n=s/2,o=Math.abs(e-n),l=s*.06,i=s*.25;let r=0;o<l&&(r=.08*(1-o/l)),o<i&&(r+=.03*(1-o/i));const $=t/h;r*=.3+$*.7;const v=1-(1-$)*.4;return o/v>i&&(r*=.2),r}function S(){d++,a.clearRect(0,0,p,x),a.font=`${c-4}px 'JetBrains Mono', monospace`,a.textBaseline="top",d>30&&Math.random()<.4&&B();for(let e=0;e<h;e++)for(let t=0;t<s;t++){const n=w[e][t];n.timer++,n.timer>n.interval&&(n.timer=0,n.char=m[Math.floor(Math.random()*m.length)]);const o=F(t,e),l=Math.random()<.002?.06:0,i=o+l;if(i<.008)continue;const r=Math.sin(d*.015+t*.1+e*.08)*.3+.7,$=Math.sin(d*.03-e*.15+t*.05)*.4+.6,v=i*r*$;if(v<.008)continue;const R=Math.floor(100+o*800),L=Math.min(220,R);a.fillStyle=`rgba(${L},${L},${L},${Math.min(.5,v)})`,a.fillText(n.char,t*M,e*c)}for(let e=g.length-1;e>=0;e--){const t=g[e];if(t.life++,t.y-=t.speed,t.x+=t.drift+Math.sin(d*.02+e)*.01,t.life%8===0&&(t.char=y[Math.floor(Math.random()*y.length)]),t.life>t.maxLife||t.y<-2){g.splice(e,1);continue}const n=t.life/t.maxLife;let o;n<.1?o=n/.1:n>.7?o=(1-n)/.3:o=1,o*=t.brightness;const l=Math.floor(160+o*90);if(a.fillStyle=`rgba(${l},${l},${l},${o*.8})`,a.fillText(t.char,Math.floor(t.x)*M,Math.floor(t.y)*c),o>.3){const i=o*.2,r=Math.floor(120+i*60);a.fillStyle=`rgba(${r},${r},${r},${i})`,a.fillText(t.char,Math.floor(t.x)*M,Math.floor(t.y+1)*c),o>.5&&(a.fillStyle=`rgba(${r},${r},${r},${i*.5})`,a.fillText(t.char,Math.floor(t.x)*M,Math.floor(t.y+2)*c))}}if(d%30===0){const e=Math.floor(s/2+(Math.random()-.5)*4),t=Math.floor(h*.4+Math.random()*h*.5);a.fillStyle="rgba(250,250,250,0.3)",a.fillText(m[Math.floor(Math.random()*m.length)],e*M,t*c)}C=requestAnimationFrame(S)}E(),S();const A=()=>E();window.addEventListener("resize",A),document.addEventListener("astro:before-swap",()=>{cancelAnimationFrame(C),window.removeEventListener("resize",A)},{once:!0})}document.readyState==="loading"?document.addEventListener("DOMContentLoaded",b):b();document.addEventListener("astro:page-load",b);</script>