<?xml version="1.0" encoding="utf-8" standalone="yes"?>
<oembed>
  <author_name>ggen-t-fukui</author_name>
  <author_url>https://blog.hatena.ne.jp/ggen-t-fukui/</author_url>
  <blog_title>G-gen Tech Blog</blog_title>
  <blog_url>https://blog.g-gen.co.jp/</blog_url>
  <categories>
    <anon>Document AI</anon>
    <anon>生成AI (Generative AI)</anon>
    <anon>Google Cloud</anon>
    <anon>フレームワーク・ガイドライン・事例</anon>
  </categories>
  <description>G-genの福井です。Google Cloud の Document AI を使い、独自 OCR モデルを開発する手順を紹介します。 はじめに 当記事の概要 Document AI とは カスタム エクストラクタとは 事前準備 サンプルレシート画像の準備 カスタム エクストラクタの作成 プロセッサの作成 ラベルの定義 データセットの準備とインポート アノテーション作業 デプロイ・テスト 基盤モデルの新しいバージョンを作成 デプロイ テスト API 呼び出し はじめに 当記事の概要 当記事では、Google Cloud が提供する Document AI のカスタム エクストラクタ機能を使用して…</description>
  <height>190</height>
  <html>&lt;iframe src=&quot;https://hatenablog-parts.com/embed?url=https%3A%2F%2Fblog.g-gen.co.jp%2Fentry%2Ftrain-custom-extractor-with-document-ai&quot; title=&quot;Document AIで独自OCRモデルをトレーニングしてみた - G-gen Tech Blog&quot; class=&quot;embed-card embed-blogcard&quot; scrolling=&quot;no&quot; frameborder=&quot;0&quot; style=&quot;display: block; width: 100%; height: 190px; max-width: 500px; margin: 10px 0px;&quot;&gt;&lt;/iframe&gt;</html>
  <image_url>https://cdn-ak.f.st-hatena.com/images/fotolife/g/ggen-sugimura/20250515/20250515185405.png</image_url>
  <provider_name>Hatena Blog</provider_name>
  <provider_url>https://hatena.blog</provider_url>
  <published>2025-05-15 09:00:00</published>
  <title>Document AIで独自OCRモデルをトレーニングしてみた</title>
  <type>rich</type>
  <url>https://blog.g-gen.co.jp/entry/train-custom-extractor-with-document-ai</url>
  <version>1.0</version>
  <width>100%</width>
</oembed>
