{"schemaVersion":"jobsearcher.job.v1","id":"ae4c5286928842b4e42d6eb8","url":"https://jobsearcher.com/jobs/ae4c5286928842b4e42d6eb8","canonicalUrl":"https://jobsearcher.com/jobs/ae4c5286928842b4e42d6eb8","title":"Performance Engineer: Optimize GPU Inference at Scale","description":"Morph is hiring a performance engineer to make the entire system faster, cheaper, and more reliable. You will work directly with the founders on problems that determine how efficiently frontier-scale models can be served, in a small team with enormous compute and immediate production impact.\nYou’ll trace latency and throughput from the API layer down to individual kernels, optimize batching, routing, quantization, and distributed execution, and build benchmarks and observability to make\n\n#J-18808-Ljbffr","company":"Techtwitterio","rawCompany":"techtwitterio","city":"Millbrae","state":"CA","isRemote":false,"isActive":false,"createdAt":"2026-09-19T04:03:19.501Z","occupations":[{"code":"15-1299.08","title":"Computer Systems Engineers/Architects","slug":"computer-systems-engineers-architects"},{"code":"15-1252.00","title":"Software Developers","slug":"software-developers"},{"code":"15-1221.00","title":"Computer and Information Research Scientists","slug":"computer-and-information-research-scientists"}],"industries":[{"code":"541511","title":"Custom Computer Programming Services","slug":"custom-computer-programming-services"},{"code":"513210","title":"Software Publishers","slug":"software-publishers"},{"code":"518210","title":"Computing Infrastructure Providers, Data Processing, Web Hosting, and Related Services","slug":"computing-infrastructure-providers-data-processing-web-hosting-and-related-services"}],"jobPosting":{"@context":"https://schema.org","@type":"JobPosting","title":"Performance Engineer: Optimize GPU Inference at Scale","description":"Morph is hiring a performance engineer to make the entire system faster, cheaper, and more reliable. You will work directly with the founders on problems that determine how efficiently frontier-scale models can be served, in a small team with enormous compute and immediate production impact.\nYou’ll trace latency and throughput from the API layer down to individual kernels, optimize batching, routing, quantization, and distributed execution, and build benchmarks and observability to make\n\n#J-18808-Ljbffr","datePosted":"2026-09-19T04:03:19.501Z","dateModified":"2026-09-19T04:03:19.501Z","hiringOrganization":{"@type":"Organization","name":"Techtwitterio","sameAs":"https://jobsearcher.com"},"jobLocation":{"@type":"Place","address":{"@type":"PostalAddress","addressLocality":"Millbrae","addressRegion":"CA","addressCountry":"US"}},"identifier":{"@type":"PropertyValue","name":"JobSearcher","value":"ae4c5286928842b4e42d6eb8"},"url":"https://jobsearcher.com/jobs/ae4c5286928842b4e42d6eb8"}}