<?xml version="1.0" encoding="UTF-8"?>
<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9">
  <url><loc>https://blog.inference-labs.com/</loc><lastmod>2026-05-29</lastmod><changefreq>daily</changefreq><priority>1.0</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/how-to-reduce-llm-api-costs-in-production</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/multi-model-llm-routing-architecture-guide</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/vendor-neutral-ai-routing-explained</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/llm-evaluation-pipeline-for-production</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/ai-inference-cost-optimization-strategies</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/when-to-use-gpt-vs-claude-vs-gemini</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/llm-fallback-and-retry-strategies-production</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/automated-ai-quality-assurance-pipeline-guide</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/llm-observability-best-practices-2026</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/ai-model-selection-policy-engineering</loc><lastmod>2026-05-27</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/reduce-cloud-ai-api-spend-startup-engineering-playbook</loc><lastmod>2026-05-28</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
  <url><loc>https://blog.inference-labs.com/posts/llm-trace-collection-production-system-design</loc><lastmod>2026-05-29</lastmod><changefreq>monthly</changefreq><priority>0.8</priority></url>
</urlset>