Happytails.aiFree tools. No account needed.

Voice to text | agent view

The same page content and links, with explicit tool capabilities. Tools marked with an API can run directly through HTTP. Other tools require the browser interface or are not connected yet.

Tool capabilities

{
  "id": "voice-to-text",
  "category": "speech-tools",
  "built": true,
  "requirements": "Record microphone dictation with explicit start and stop; distinguish provisional from final words; edit and copy text; disclose local versus uploaded recognition before recording.",
  "implementation": {
    "fields": [],
    "files": true,
    "accept": ".wav,.mp3,.m4a,.flac,.mp4,.mov,.webm",
    "multiple": false,
    "input": false,
    "processing": "Server required",
    "native": true,
    "ai": true,
    "maxFileBytes": 20971520,
    "note": "Runs with local model weights on this website’s server. No external AI API is used. Review the result for errors. Local Whisper tiny. Clips up to five minutes; review words and timing. Upload a voice recording; direct microphone capture is not yet included.",
    "experience": {
      "automatic": false,
      "single": false,
      "label": "Your text",
      "placeholder": "Type or paste your input…",
      "help": "Up to 1,000,000 characters. Review the settings before running.",
      "action": "Convert file"
    },
    "mode": "explicit",
    "interactive": false
  },
  "execution": "POST /api/v1/tools/voice-to-text/run",
  "api": {
    "id": "voice-to-text",
    "name": "Voice to text",
    "category": "speech-tools",
    "description": "Turn an uploaded voice recording into a transcript.",
    "notes": "Runs with local model weights on this website’s server. No external AI API is used. Review the result for errors. Local Whisper tiny. Clips up to five minutes; review words and timing. Upload a voice recording; direct microphone capture is not yet included.",
    "human_path": "/en/speech-tools/voice-to-text/",
    "method": "POST",
    "endpoint": "/api/v1/tools/voice-to-text/run",
    "schema_url": "/api/v1/tools/voice-to-text",
    "input_schema": {
      "type": "object",
      "additionalProperties": false,
      "properties": {
        "input": {
          "type": "string",
          "maxLength": 1000000,
          "default": "",
          "description": "Plain text input. File tools use files instead unless otherwise documented."
        },
        "options": {
          "type": "object",
          "additionalProperties": false,
          "properties": {}
        },
        "files": {
          "type": "array",
          "maxItems": 1,
          "items": {
            "type": "object",
            "required": [
              "name",
              "base64"
            ],
            "additionalProperties": false,
            "properties": {
              "name": {
                "type": "string",
                "maxLength": 200,
                "description": "Filename only, no path."
              },
              "mime": {
                "type": "string",
                "maxLength": 150
              },
              "bytes": {
                "type": "integer",
                "minimum": 0,
                "description": "Optional decoded byte count; must match content if supplied."
              },
              "base64": {
                "type": "string",
                "contentEncoding": "base64",
                "description": "File bytes as padded base64. Use the tool-specific limits.files_bytes value for the total decoded input size."
              }
            }
          }
        }
      }
    },
    "output_schema": {
      "type": "object",
      "required": [
        "tool",
        "result"
      ],
      "properties": {
        "tool": {
          "type": "string"
        },
        "result": {
          "type": "object",
          "required": [
            "text",
            "files"
          ],
          "properties": {
            "text": {
              "type": [
                "string",
                "null"
              ]
            },
            "data": {
              "description": "Parsed JSON when the textual result is JSON."
            },
            "name": {
              "type": "string"
            },
            "mime": {
              "type": "string"
            },
            "files": {
              "type": "array",
              "items": {
                "type": "object",
                "required": [
                  "name",
                  "mime",
                  "bytes",
                  "base64"
                ],
                "properties": {
                  "name": {
                    "type": "string"
                  },
                  "mime": {
                    "type": "string"
                  },
                  "bytes": {
                    "type": "integer"
                  },
                  "base64": {
                    "type": "string",
                    "contentEncoding": "base64"
                  }
                }
              }
            }
          }
        }
      }
    },
    "processing": "Server; self-hosted native engine, no external conversion service",
    "limits": {
      "request_bytes": 31457280,
      "input_characters": 1000000,
      "files_bytes": 20971520,
      "max_files": 1,
      "timeout_seconds": 90
    }
  },
  "agent_processing": "Server; self-hosted native engine, no external conversion service",
  "browser_agent": {
    "supported": true,
    "name": "run_current_tool",
    "discovery": "WebMCP on the human page in a supporting browser",
    "verification": "See audit/API-AUDIT.md; availability is not verification"
  }
}

Page guide

Turn an uploaded voice recording into a transcript.

Useful next steps

How to use Voice to text

  1. Choose a file in WAV, MP3, M4A, FLAC, MP4, MOV or WEBM.
  2. Review the settings, then choose Convert file.
  3. Review the result, then use Copy result or a download link when available.

What to expect

Turn an uploaded voice recording into a transcript. Runs with local model weights on this website’s server. No external AI API is used. Review the result for errors. Local Whisper tiny. Clips up to five minutes; review words and timing. Upload a voice recording; direct microphone capture is not yet included.