Blame view

src/lm/arpa-lm-compiler.h 1.84 KB
8dcb6dfcb   Yannick Estève   first commit
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
  // lm/arpa-lm-compiler.h
  
  // Copyright 2009-2011 Gilles Boulianne
  // Copyright 2016 Smart Action LLC (kkm)
  
  // See ../../COPYING for clarification regarding multiple authors
  //
  // Licensed under the Apache License, Version 2.0 (the "License");
  // you may not use this file except in compliance with the License.
  // You may obtain a copy of the License at
  //
  //  http://www.apache.org/licenses/LICENSE-2.0
  //
  // THIS CODE IS PROVIDED *AS IS* BASIS, WITHOUT WARRANTIES OR CONDITIONS OF ANY
  // KIND, EITHER EXPRESS OR IMPLIED, INCLUDING WITHOUT LIMITATION ANY IMPLIED
  // WARRANTIES OR CONDITIONS OF TITLE, FITNESS FOR A PARTICULAR PURPOSE,
  // MERCHANTABLITY OR NON-INFRINGEMENT.
  // See the Apache 2 License for the specific language governing permissions and
  // limitations under the License.
  
  #ifndef KALDI_LM_ARPA_LM_COMPILER_H_
  #define KALDI_LM_ARPA_LM_COMPILER_H_
  
  #include <fst/fstlib.h>
  
  #include "lm/arpa-file-parser.h"
  
  namespace kaldi {
  
  class ArpaLmCompilerImplInterface;
  
  class ArpaLmCompiler : public ArpaFileParser {
   public:
    ArpaLmCompiler(ArpaParseOptions options, int sub_eps,
                   fst::SymbolTable* symbols)
        : ArpaFileParser(options, symbols),
          sub_eps_(sub_eps), impl_(NULL) {
    }
    ~ArpaLmCompiler();
  
    const fst::StdVectorFst& Fst() const { return fst_; }
    fst::StdVectorFst* MutableFst() { return &fst_; }
  
   protected:
    // ArpaFileParser overrides.
    virtual void HeaderAvailable();
    virtual void ConsumeNGram(const NGram& ngram);
    virtual void ReadComplete();
  
  
   private:
    // this function removes states that only have a backoff arc coming
    // out of them.
    void RemoveRedundantStates();
    void Check() const;
  
    int sub_eps_;
    ArpaLmCompilerImplInterface* impl_;  // Owned.
    fst::StdVectorFst fst_;
    template <class HistKey> friend class ArpaLmCompilerImpl;
  };
  
  }  // namespace kaldi
  
  #endif  // KALDI_LM_ARPA_LM_COMPILER_H_