Happytails.aiFree tools. No account needed.

Audio to text | agent view

The same page content and links, with explicit tool capabilities. Tools marked with an API can run directly through HTTP. Other tools require the browser interface or are not connected yet.

Tool capabilities

{
  "id": "audio-to-text",
  "category": "speech-tools",
  "built": true,
  "requirements": "Upload audio; select or detect language; edit timestamped transcript; export TXT and JSON; show uncertain passages and suppress silence hallucinations. Test accents, noise and long-file chunk boundaries.",
  "implementation": {
    "fields": [],
    "files": true,
    "accept": ".wav,.mp3,.m4a,.flac,.mp4,.mov,.webm",
    "multiple": false,
    "input": false,
    "processing": "Server required",
    "native": true,
    "ai": true,
    "maxFileBytes": 20971520,
    "note": "Runs with local model weights on this website’s server. No external AI API is used. Review the result for errors. Local Whisper tiny. Clips up to five minutes; review words and timing.",
    "experience": {
      "automatic": false,
      "single": false,
      "label": "Your text",
      "placeholder": "Type or paste your input…",
      "help": "Up to 1,000,000 characters. Review the settings before running.",
      "action": "Convert file"
    },
    "mode": "explicit",
    "interactive": false
  },
  "execution": "POST /api/v1/tools/audio-to-text/run",
  "api": {
    "id": "audio-to-text",
    "name": "Audio to text",
    "category": "speech-tools",
    "description": "Transcribe a short audio recording into text and timed JSON.",
    "notes": "Runs with local model weights on this website’s server. No external AI API is used. Review the result for errors. Local Whisper tiny. Clips up to five minutes; review words and timing.",
    "human_path": "/en/speech-tools/audio-to-text/",
    "method": "POST",
    "endpoint": "/api/v1/tools/audio-to-text/run",
    "schema_url": "/api/v1/tools/audio-to-text",
    "input_schema": {
      "type": "object",
      "additionalProperties": false,
      "properties": {
        "input": {
          "type": "string",
          "maxLength": 1000000,
          "default": "",
          "description": "Plain text input. File tools use files instead unless otherwise documented."
        },
        "options": {
          "type": "object",
          "additionalProperties": false,
          "properties": {}
        },
        "files": {
          "type": "array",
          "maxItems": 1,
          "items": {
            "type": "object",
            "required": [
              "name",
              "base64"
            ],
            "additionalProperties": false,
            "properties": {
              "name": {
                "type": "string",
                "maxLength": 200,
                "description": "Filename only, no path."
              },
              "mime": {
                "type": "string",
                "maxLength": 150
              },
              "bytes": {
                "type": "integer",
                "minimum": 0,
                "description": "Optional decoded byte count; must match content if supplied."
              },
              "base64": {
                "type": "string",
                "contentEncoding": "base64",
                "description": "File bytes as padded base64. Use the tool-specific limits.files_bytes value for the total decoded input size."
              }
            }
          }
        }
      }
    },
    "output_schema": {
      "type": "object",
      "required": [
        "tool",
        "result"
      ],
      "properties": {
        "tool": {
          "type": "string"
        },
        "result": {
          "type": "object",
          "required": [
            "text",
            "files"
          ],
          "properties": {
            "text": {
              "type": [
                "string",
                "null"
              ]
            },
            "data": {
              "description": "Parsed JSON when the textual result is JSON."
            },
            "name": {
              "type": "string"
            },
            "mime": {
              "type": "string"
            },
            "files": {
              "type": "array",
              "items": {
                "type": "object",
                "required": [
                  "name",
                  "mime",
                  "bytes",
                  "base64"
                ],
                "properties": {
                  "name": {
                    "type": "string"
                  },
                  "mime": {
                    "type": "string"
                  },
                  "bytes": {
                    "type": "integer"
                  },
                  "base64": {
                    "type": "string",
                    "contentEncoding": "base64"
                  }
                }
              }
            }
          }
        }
      }
    },
    "processing": "Server; self-hosted native engine, no external conversion service",
    "limits": {
      "request_bytes": 31457280,
      "input_characters": 1000000,
      "files_bytes": 20971520,
      "max_files": 1,
      "timeout_seconds": 90
    }
  },
  "agent_processing": "Server; self-hosted native engine, no external conversion service",
  "browser_agent": {
    "supported": true,
    "name": "run_current_tool",
    "discovery": "WebMCP on the human page in a supporting browser",
    "verification": "See audit/API-AUDIT.md; availability is not verification"
  }
}

Page guide

Transcribe a short audio recording into text and timed JSON.

Useful next steps

How to use Audio to text

  1. Choose a file in WAV, MP3, M4A, FLAC, MP4, MOV or WEBM.
  2. Review the settings, then choose Convert file.
  3. Review the result, then use Copy result or a download link when available.

What to expect

Transcribe a short audio recording into text and timed JSON. Runs with local model weights on this website’s server. No external AI API is used. Review the result for errors. Local Whisper tiny. Clips up to five minutes; review words and timing.