{"origin":{"source":"Infrabase.ai","url":"https://infrabase.ai","license":"CC-BY-4.0","note":"Directory data maintained by hand at infrabase.ai. Attribution required.","docs":"https://infrabase.ai/api/query/docs"},"count":20,"results":[{"name":"AiQu","slug":"aiqu","infrabase_url":"https://infrabase.ai/inference-apis/aiqu","site_url":"https://aiqu.ai","tagline":"Swedish GPU infrastructure and LLM hosting platform with API-first deployment, no Kubernetes required","primary_job":"gpu-compute","secondary_jobs":["hosted-inference-api"],"categories":["inference-apis"],"pricing":"Monthly subscriptions","hq_country":"SE","gdpr":true,"github_stars":null},{"name":"Airon","slug":"airon","infrabase_url":"https://infrabase.ai/inference-apis/airon","site_url":"https://airon.ai","tagline":"Dedicated bare-metal GPU infrastructure for AI workloads, hosted in Nordic datacenters","primary_job":"gpu-compute","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Hourly","hq_country":"SE","gdpr":true,"github_stars":null},{"name":"AKI.IO","slug":"aki-io","infrabase_url":"https://infrabase.ai/inference-apis/aki-io","site_url":"https://aki.io","tagline":"European AI API for open-source models on EU infrastructure","primary_job":"gpu-compute","secondary_jobs":["hosted-inference-api"],"categories":["inference-apis"],"pricing":"10 EUR free credits, then pay-per-token","hq_country":"DE","gdpr":null,"github_stars":null},{"name":"Amazon Bedrock","slug":"amazon-bedrock","infrabase_url":"https://infrabase.ai/inference-apis/amazon-bedrock","site_url":"https://aws.amazon.com/bedrock/","tagline":"Managed API access to foundation models on AWS with built-in fine-tuning and agent tooling","primary_job":"hosted-inference-api","secondary_jobs":["fine-tuning-platform"],"categories":["inference-apis"],"pricing":"Per token usage","hq_country":"US","gdpr":true,"github_stars":null},{"name":"Anthropic Claude","slug":"anthropic-claude","infrabase_url":"https://infrabase.ai/inference-apis/anthropic-claude","site_url":"https://claude.ai/","tagline":"Claude API for building AI applications with Opus, Sonnet, and Haiku models","primary_job":"hosted-inference-api","secondary_jobs":["agent-platform"],"categories":["inference-apis"],"pricing":"Per token usage","hq_country":"US","gdpr":true,"github_stars":3523},{"name":"Anyscale","slug":"anyscale","infrabase_url":"https://infrabase.ai/inference-apis/anyscale","site_url":"https://www.anyscale.com/","tagline":"Fast, cost-efficient, serverless APIs for LLM Serving and Fine Tuning","primary_job":"hosted-inference-api","secondary_jobs":["fine-tuning-platform","serverless-gpu-platform"],"categories":["inference-apis"],"pricing":"Pay-as-you-go","hq_country":"US","gdpr":null,"github_stars":null},{"name":"ARK Labs","slug":"ark-labs","infrabase_url":"https://infrabase.ai/inference-apis/ark-labs","site_url":"https://ark-labs.cloud","tagline":"Sovereign AI inference infrastructure for regulated EU environments, with heterogeneous GPU support","primary_job":"gpu-compute","secondary_jobs":["hosted-inference-api"],"categories":["inference-apis"],"pricing":"Pay-as-you-go","hq_country":null,"gdpr":null,"github_stars":null},{"name":"Baseten","slug":"baseten","infrabase_url":"https://infrabase.ai/inference-apis/baseten","site_url":"https://www.baseten.co","tagline":"AI inference platform for deploying and serving ML models with autoscaling and optimized infrastructure","primary_job":"serverless-gpu-platform","secondary_jobs":["hosted-inference-api"],"categories":["inference-apis"],"pricing":"Usage-based","hq_country":"US","gdpr":null,"github_stars":1155},{"name":"Beam","slug":"beam","infrabase_url":"https://infrabase.ai/inference-apis/beam","site_url":"https://www.beam.cloud","tagline":"Open-source serverless GPU cloud with sub-second cold starts and auto-scaling","primary_job":"serverless-gpu-platform","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Pay-as-you-go","hq_country":"US","gdpr":null,"github_stars":1654},{"name":"BentoML","slug":"bentoml","infrabase_url":"https://infrabase.ai/inference-apis/bentoml","site_url":"https://bentoml.com/","tagline":"BentoML is the platform for software engineers to build AI products.","primary_job":"serverless-gpu-platform","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Pay-as-you-go","hq_country":"US","gdpr":null,"github_stars":null},{"name":"Berget AI","slug":"berget-ai","infrabase_url":"https://infrabase.ai/inference-apis/berget-ai","site_url":"https://berget.ai","tagline":"EU-sovereign AI inference platform with OpenAI-compatible API","primary_job":"hosted-inference-api","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Free tier with €5 credits. Starter €25/mo, Team €50/mo, Enterprise €500/mo.","hq_country":"SE","gdpr":true,"github_stars":1},{"name":"Cerebras","slug":"cerebras","infrabase_url":"https://infrabase.ai/inference-apis/cerebras","site_url":"https://cerebras.ai","tagline":"Ultra-fast inference on custom wafer-scale hardware with OpenAI-compatible API","primary_job":"hosted-inference-api","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Per token usage","hq_country":"US","gdpr":null,"github_stars":1165},{"name":"Cerebrium","slug":"cerebrium","infrabase_url":"https://infrabase.ai/inference-apis/cerebrium","site_url":"https://www.cerebrium.ai","tagline":"Serverless GPU infrastructure for deploying AI models with sub-5 second cold starts","primary_job":"serverless-gpu-platform","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Pay-per-second","hq_country":"US","gdpr":null,"github_stars":null},{"name":"CheapestInference","slug":"cheapestinference","infrabase_url":"https://infrabase.ai/inference-apis/cheapestinference","site_url":"https://cheapestinference.com","tagline":"Flat-rate unlimited inference on open-weight models, sold in daily 8-hour windows","primary_job":"hosted-inference-api","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Monthly subscriptions","hq_country":null,"gdpr":null,"github_stars":null},{"name":"Cloudflare Workers AI","slug":"cloudflare-workers-ai","infrabase_url":"https://infrabase.ai/inference-apis/cloudflare-workers-ai","site_url":"https://developers.cloudflare.com/workers-ai/","tagline":"Run AI models at the edge on Cloudflare's global network with serverless inference","primary_job":"hosted-inference-api","secondary_jobs":["serverless-gpu-platform"],"categories":["inference-apis"],"pricing":"Usage-based","hq_country":"US","gdpr":null,"github_stars":null},{"name":"CodingPlanX","slug":"codingplanx","infrabase_url":"https://infrabase.ai/inference-apis/codingplanx","site_url":"https://codingplanx.ai","tagline":"Unified AI API gateway providing access to 600+ models from OpenAI, Anthropic, Google, DeepSeek, and more","primary_job":"hosted-inference-api","secondary_jobs":["llm-gateway"],"categories":["inference-apis"],"pricing":"Usage-based","hq_country":null,"gdpr":null,"github_stars":null},{"name":"cohere","slug":"cohere","infrabase_url":"https://infrabase.ai/inference-apis/cohere","site_url":"https://cohere.com/","tagline":"Cohere’s world-class LLMs help enterprises build powerful, secure applications that search, understand meaning and converse in text.","primary_job":"hosted-inference-api","secondary_jobs":["embedding-api"],"categories":["inference-apis"],"pricing":"Per token usage","hq_country":"CA","gdpr":true,"github_stars":null},{"name":"CoreWeave","slug":"coreweave","infrabase_url":"https://infrabase.ai/inference-apis/coreweave","site_url":"https://www.coreweave.com","tagline":"GPU cloud infrastructure built for large-scale AI training and inference workloads","primary_job":"gpu-compute","secondary_jobs":[],"categories":["inference-apis"],"pricing":"Usage-based","hq_country":"US","gdpr":null,"github_stars":null},{"name":"Cortecs AI","slug":"cortecs-ai","infrabase_url":"https://infrabase.ai/inference-apis/cortecs-ai","site_url":"https://cortecs.ai","tagline":"European AI inference gateway with smart routing across EU providers","primary_job":"hosted-inference-api","secondary_jobs":["llm-gateway"],"categories":["inference-apis"],"pricing":null,"hq_country":"AT","gdpr":true,"github_stars":null},{"name":"deepinfra","slug":"deepinfra","infrabase_url":"https://infrabase.ai/inference-apis/deepinfra","site_url":"https://deepinfra.com/","tagline":"Run the top AI models using a simple API, pay per use. Low cost, scalable and production ready infrastructure.","primary_job":"hosted-inference-api","secondary_jobs":["serverless-gpu-platform"],"categories":["inference-apis"],"pricing":"Per token usage","hq_country":"US","gdpr":true,"github_stars":null}]}