Module Values.DatasetEntityRecognizerDocumentsSource

Describes the documents submitted with a dataset for an entity recognizer model.

Sourcetype nonrec t = {
  1. s3Uri : S3Uri.t;
    (*

    Specifies the Amazon S3 location where the documents for the dataset are located.

    *)
  2. inputFormat : InputFormat.t option;
    (*

    Specifies how the text in an input file should be processed. This is optional, and the default is ONE_DOC_PER_LINE. ONE_DOC_PER_FILE - Each file is considered a separate document. Use this option when you are processing large documents, such as newspaper articles or scientific papers. ONE_DOC_PER_LINE - Each line in a file is considered a separate document. Use this option when you are processing many short documents, such as text messages.

    *)
}
Sourceval context_ : string
Sourceval make : ?inputFormat:??? -> s3Uri:S3Uri.t -> unit -> t
Sourceval to_value : t -> [> `Structure of (string * [> `Enum of string | `String of S3Uri.t ]) list ]
Sourceval to_query : t -> Awso.Client.Query.t
Sourceval of_xml : Awso.Xml.t -> t
Sourceval of_string : string -> t
Sourceval of_json : Yojson.Safe.t -> t
Sourceval to_json : t -> Yojson.Safe.t