{"schemaVersion":"jobsearcher.job.v1","id":"dad31f26b5ddc1ed80d9a8fb","url":"https://jobsearcher.com/jobs/dad31f26b5ddc1ed80d9a8fb","canonicalUrl":"https://jobsearcher.com/jobs/dad31f26b5ddc1ed80d9a8fb","title":"Language Model Evaluator - Fully Remote","description":"Role ResponsibilitiesConduct fact-checking using trusted public sources and externaltools .Generate high-quality human evaluation data by identifying response strengths, areas for improvement, and factual inaccuracies.Assess reasoning quality, clarity, tone, and completeness of responses.Ensure model responses align with expected conversational behavior and system guidelines.Workindependently and asynchronouslyto meet deadlines while improvingAI model performance .Qualifications\r\nMust-HaveExcellent writing skills in EnglishStrong attention to detailBackground or experience in domains requiringstructured analytical thinking(e.g., research, policy, analytics, linguistics, engineering)PreferredPrior experience withRLHF, model evaluation, or data annotation workExperience writing or editinghigh-quality written contentExperience comparing multiple outputs and makingfine-grained qualitative judgments#J-18808-Ljbffr","company":"Mercor","rawCompany":"mercor","city":"Millbrae","state":"CA","isRemote":true,"isActive":false,"createdAt":"2026-07-03T01:11:46.425Z","occupations":[{"code":"27-3043.00","title":"Writers and Authors","slug":"writers-and-authors"},{"code":"27-3091.00","title":"Interpreters and Translators","slug":"interpreters-and-translators"},{"code":"27-3041.00","title":"Editors","slug":"editors"}],"industries":[{"code":"541930","title":"Translation and Interpretation Services","slug":"translation-and-interpretation-services"},{"code":"541990","title":"All Other Professional, Scientific, and Technical Services","slug":"all-other-professional-scientific-and-technical-services"},{"code":"541690","title":"Other Scientific and Technical Consulting Services","slug":"other-scientific-and-technical-consulting-services"}],"jobPosting":{"@context":"https://schema.org","@type":"JobPosting","title":"Language Model Evaluator - Fully Remote","description":"Role ResponsibilitiesConduct fact-checking using trusted public sources and externaltools .Generate high-quality human evaluation data by identifying response strengths, areas for improvement, and factual inaccuracies.Assess reasoning quality, clarity, tone, and completeness of responses.Ensure model responses align with expected conversational behavior and system guidelines.Workindependently and asynchronouslyto meet deadlines while improvingAI model performance .Qualifications\r\nMust-HaveExcellent writing skills in EnglishStrong attention to detailBackground or experience in domains requiringstructured analytical thinking(e.g., research, policy, analytics, linguistics, engineering)PreferredPrior experience withRLHF, model evaluation, or data annotation workExperience writing or editinghigh-quality written contentExperience comparing multiple outputs and makingfine-grained qualitative judgments#J-18808-Ljbffr","datePosted":"2026-07-03T01:11:46.425Z","dateModified":"2026-07-03T01:11:46.425Z","hiringOrganization":{"@type":"Organization","name":"Mercor","sameAs":"https://jobsearcher.com"},"jobLocationType":"TELECOMMUTE","applicantLocationRequirements":{"@type":"Country","name":"US"},"jobLocation":{"@type":"Place","address":{"@type":"PostalAddress","addressLocality":"Millbrae","addressRegion":"CA","addressCountry":"US"}},"identifier":{"@type":"PropertyValue","name":"JobSearcher","value":"dad31f26b5ddc1ed80d9a8fb"},"url":"https://jobsearcher.com/jobs/dad31f26b5ddc1ed80d9a8fb"}}